new-model / tokenizer_config.json
Efe2898's picture
Add RSLM tokenizer trained on AkademikDerlem/makaleler
1966da3 verified
raw
history blame contribute delete
818 Bytes
{
"model_max_length": 262144,
"tokenizer_class": "PreTrainedTokenizerFast",
"clean_up_tokenization_spaces": false,
"padding_side": "right",
"truncation_side": "right",
"bos_token": "<|bos|>",
"eos_token": "<|eos|>",
"unk_token": "<|unk|>",
"pad_token": "<|pad|>",
"additional_special_tokens": [
"<|system|>",
"<|user|>",
"<|assistant|>",
"<|answer|>",
"<|end|>",
"<think>",
"</think>"
],
"chat_template": "{% for message in messages %}{% if message['role'] == 'system' %}<|system|>\n{{ message['content'] }}<|end|>\n{% elif message['role'] == 'user' %}<|user|>\n{{ message['content'] }}<|end|>\n{% elif message['role'] == 'assistant' %}<|assistant|>\n{{ message['content'] }}<|end|>\n{% endif %}{% endfor %}{% if add_generation_prompt %}<|assistant|>\n{% endif %}"
}