testai / tokenizer_config.json
testaccount-ai's picture
Create tokenizer_config.json
17c19b0 verified
{
"add_bos_token": false,
"add_eos_token": false,
"add_prefix_space": true,
"added_tokens_decoder": {
"151643": {
"content": "<|endoftext|>",
"lstrip": false,
"normalized": false,
"rstrip": false,
"single_word": false,
"special": true
},
"151644": {
"content": "<|user|>",
"lstrip": false,
"normalized": false,
"rstrip": false,
"single_word": false,
"special": true
},
"151645": {
"content": "<|assistant|>",
"lstrip": false,
"normalized": false,
"rstrip": false,
"single_word": false,
"special": true
},
"151646": {
"content": "<|system|>",
"lstrip": false,
"normalized": false,
"rstrip": false,
"single_word": false,
"special": true
}
},
"bos_token": "<|endoftext|>",
"chat_template": "{%- set ns = namespace(messages=[], system_prompt='') -%}{%- for message in messages -%}{%- if message['role'] == 'system' -%}{%- set ns.system_prompt = message['content'] -%}{%- else -%}{%- set ns.messages = ns.messages + [message] -%}{%- endif -%}{%- endfor -%}{%- if ns.system_prompt != '' -%}{%- set system_message = '<|system|>\\n' + ns.system_prompt + '\\n' -%}{%- else -%}{%- set system_message = '' -%}{%- endif -%}{%- for message in ns.messages -%}{%- if message['role'] == 'user' -%}{{ '<|user|>\\n' + message['content'] + '\\n' }}{%- elif message['role'] == 'assistant' -%}{{ '<|assistant|>\\n' + message['content'] + '\\n' }}{%- endif -%}{%- endfor -%}{%- if add_generation_prompt -%}{{ '<|assistant|>\\n' }}{%- endif -%}",
"clean_up_tokenization_spaces": false,
"eos_token": "<|endoftext|>",
"legacy": false,
"model_max_length": 8192,
"pad_token": "<|endoftext|>",
"padding_side": "right",
"sp_model_kwargs": {},
"tokenizer_class": "PreTrainedTokenizerFast",
"trust_remote_code": true,
"unk_token": null,
"use_fast": true
}