Flansma commited on
Commit
522bb50
·
verified ·
1 Parent(s): 62249d1

Update to best checkpoint (epoch=85, val_loss=0.3335)

Browse files
config.json CHANGED
@@ -2,6 +2,7 @@
2
  "architectures": [
3
  "HELMBertForMaskedLM"
4
  ],
 
5
  "auto_map": {
6
  "AutoConfig": "configuration_helmbert.HELMBertConfig",
7
  "AutoModel": "modeling_helmbert.HELMBertModel",
@@ -9,8 +10,9 @@
9
  "AutoModelForSequenceClassification": "modeling_helmbert.HELMBertForSequenceClassification",
10
  "AutoTokenizer": "tokenization_helmbert.HELMBertTokenizer"
11
  },
12
- "attention_probs_dropout_prob": 0.1,
13
  "bos_token_id": 1,
 
 
14
  "dtype": "float32",
15
  "eos_token_id": 2,
16
  "hidden_dropout_prob": 0.1,
 
2
  "architectures": [
3
  "HELMBertForMaskedLM"
4
  ],
5
+ "attention_probs_dropout_prob": 0.1,
6
  "auto_map": {
7
  "AutoConfig": "configuration_helmbert.HELMBertConfig",
8
  "AutoModel": "modeling_helmbert.HELMBertModel",
 
10
  "AutoModelForSequenceClassification": "modeling_helmbert.HELMBertForSequenceClassification",
11
  "AutoTokenizer": "tokenization_helmbert.HELMBertTokenizer"
12
  },
 
13
  "bos_token_id": 1,
14
+ "classifier_dropout": 0.1,
15
+ "classifier_num_layers": 0,
16
  "dtype": "float32",
17
  "eos_token_id": 2,
18
  "hidden_dropout_prob": 0.1,
model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:c735315218d2e8864e4ea93f2c76d8eef58f78989f007b6a2bf446af6657a194
3
  size 219166144
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:6e05d977e54557fd1703afc5c053dd625bb1473621df3fc913a3c8547826b9fe
3
  size 219166144
special_tokens_map.json CHANGED
@@ -1,9 +1,51 @@
1
  {
2
- "bos_token": "@",
3
- "cls_token": "@",
4
- "eos_token": "\n",
5
- "mask_token": "¶",
6
- "pad_token": " ",
7
- "sep_token": "\n",
8
- "unk_token": "§"
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
9
  }
 
1
  {
2
+ "bos_token": {
3
+ "content": "@",
4
+ "lstrip": false,
5
+ "normalized": false,
6
+ "rstrip": false,
7
+ "single_word": false
8
+ },
9
+ "cls_token": {
10
+ "content": "@",
11
+ "lstrip": false,
12
+ "normalized": false,
13
+ "rstrip": false,
14
+ "single_word": false
15
+ },
16
+ "eos_token": {
17
+ "content": "\n",
18
+ "lstrip": false,
19
+ "normalized": false,
20
+ "rstrip": false,
21
+ "single_word": false
22
+ },
23
+ "mask_token": {
24
+ "content": "¶",
25
+ "lstrip": false,
26
+ "normalized": false,
27
+ "rstrip": false,
28
+ "single_word": false
29
+ },
30
+ "pad_token": {
31
+ "content": " ",
32
+ "lstrip": false,
33
+ "normalized": false,
34
+ "rstrip": false,
35
+ "single_word": false
36
+ },
37
+ "sep_token": {
38
+ "content": "\n",
39
+ "lstrip": false,
40
+ "normalized": false,
41
+ "rstrip": false,
42
+ "single_word": false
43
+ },
44
+ "unk_token": {
45
+ "content": "§",
46
+ "lstrip": false,
47
+ "normalized": false,
48
+ "rstrip": false,
49
+ "single_word": false
50
+ }
51
  }
tokenizer_config.json CHANGED
@@ -1,7 +1,4 @@
1
  {
2
- "auto_map": {
3
- "AutoTokenizer": ["tokenization_helmbert.HELMBertTokenizer", null]
4
- },
5
  "added_tokens_decoder": {
6
  "0": {
7
  "content": " ",
@@ -44,6 +41,12 @@
44
  "special": true
45
  }
46
  },
 
 
 
 
 
 
47
  "bos_token": "@",
48
  "clean_up_tokenization_spaces": false,
49
  "cls_token": "@",
 
1
  {
 
 
 
2
  "added_tokens_decoder": {
3
  "0": {
4
  "content": " ",
 
41
  "special": true
42
  }
43
  },
44
+ "auto_map": {
45
+ "AutoTokenizer": [
46
+ "tokenization_helmbert.HELMBertTokenizer",
47
+ null
48
+ ]
49
+ },
50
  "bos_token": "@",
51
  "clean_up_tokenization_spaces": false,
52
  "cls_token": "@",