gec_sft_v2 / tokenizer_config.json
hyan's picture
Upload folder using huggingface_hub
c6d2233 verified
{
"add_prefix_space": false,
"backend": "tokenizers",
"bos_token": "<|im_start|>",
"clean_up_tokenization_spaces": false,
"eos_token": "<|im_end|>",
"extra_special_tokens": [
"<|im_start|>",
"<|im_end|>"
],
"is_local": false,
"model_max_length": 2048,
"pad_token": "<empty_output>",
"padding_side": "left",
"tokenizer_class": "TokenizersBackend",
"unk_token": "<|endoftext|>",
"vocab_size": 49152
}