ssdatar commited on
Commit
689de6a
·
1 Parent(s): 53ef11c

Training in progress, step 5500

Browse files
config.json CHANGED
Binary files a/config.json and b/config.json differ
 
pytorch_model.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:8df328ad546730da7f824bd340345b72314f319e2a435ec2bf3eefc104b80932
3
  size 384469069
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2282b320bb3283a1fb6e978e79a88793f78a20a39caa40cfe9344146e88ba265
3
  size 384469069
special_tokens_map.json CHANGED
@@ -1,15 +1,7 @@
1
  {
2
- "bos_token": "<s>",
3
- "cls_token": "<s>",
4
- "eos_token": "</s>",
5
- "mask_token": {
6
- "content": "<mask>",
7
- "lstrip": true,
8
- "normalized": false,
9
- "rstrip": false,
10
- "single_word": false
11
- },
12
- "pad_token": "<pad>",
13
- "sep_token": "</s>",
14
- "unk_token": "<unk>"
15
  }
 
1
  {
2
+ "cls_token": "[CLS]",
3
+ "mask_token": "[MASK]",
4
+ "pad_token": "[PAD]",
5
+ "sep_token": "[SEP]",
6
+ "unk_token": "[UNK]"
 
 
 
 
 
 
 
 
7
  }
tokenizer_config.json CHANGED
@@ -1,15 +1,13 @@
1
  {
2
- "add_prefix_space": false,
3
- "bos_token": "<s>",
4
  "clean_up_tokenization_spaces": true,
5
- "cls_token": "<s>",
6
- "eos_token": "</s>",
7
- "errors": "replace",
8
- "mask_token": "<mask>",
9
- "model_max_length": 1024,
10
- "pad_token": "<pad>",
11
- "sep_token": "</s>",
12
- "tokenizer_class": "BartTokenizer",
13
- "trim_offsets": true,
14
- "unk_token": "<unk>"
15
  }
 
1
  {
 
 
2
  "clean_up_tokenization_spaces": true,
3
+ "cls_token": "[CLS]",
4
+ "do_lower_case": true,
5
+ "mask_token": "[MASK]",
6
+ "model_max_length": 512,
7
+ "pad_token": "[PAD]",
8
+ "sep_token": "[SEP]",
9
+ "strip_accents": null,
10
+ "tokenize_chinese_chars": true,
11
+ "tokenizer_class": "ElectraTokenizer",
12
+ "unk_token": "[UNK]"
13
  }