Fares11 commited on
Commit
858e42e
·
verified ·
1 Parent(s): 962f077

Upload Tokenizer

Browse files
.gitattributes CHANGED
@@ -33,3 +33,4 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
 
 
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
36
+ tokenizer.json filter=lfs diff=lfs merge=lfs -text
chat_template.jinja ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ {% for message in messages %}{{'<|im_start|>' + message['role'] + '
2
+ ' + message['content'] + '<|im_end|>' + '
3
+ '}}{% endfor %}{% if add_generation_prompt %}{{ '<|im_start|>assistant
4
+ ' }}{% endif %}
tokenizer.json ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7fad9b5f6f930b43d292eb3c56c176a69292850ddd0abc02d9ea1dac3292c87a
3
+ size 33442428
tokenizer_config.json ADDED
@@ -0,0 +1,30 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "audio_token": "<audio_soft_token>",
3
+ "backend": "tokenizers",
4
+ "boa_token": "<start_of_audio>",
5
+ "boi_token": "<start_of_image>",
6
+ "bos_token": "<bos>",
7
+ "clean_up_tokenization_spaces": false,
8
+ "eoa_token": "<end_of_audio>",
9
+ "eoi_token": "<end_of_image>",
10
+ "eos_token": "<eos>",
11
+ "image_token": "<image_soft_token>",
12
+ "is_local": false,
13
+ "mask_token": "<mask>",
14
+ "model_max_length": 1000000000000000019884624838656,
15
+ "model_specific_special_tokens": {
16
+ "audio_token": "<audio_soft_token>",
17
+ "boa_token": "<start_of_audio>",
18
+ "boi_token": "<start_of_image>",
19
+ "eoa_token": "<end_of_audio>",
20
+ "eoi_token": "<end_of_image>",
21
+ "image_token": "<image_soft_token>"
22
+ },
23
+ "pad_token": "<pad>",
24
+ "processor_class": "Gemma3nProcessor",
25
+ "sp_model_kwargs": null,
26
+ "spaces_between_special_tokens": false,
27
+ "tokenizer_class": "GemmaTokenizer",
28
+ "unk_token": "<unk>",
29
+ "use_default_system_prompt": false
30
+ }