RLau33 commited on
Commit
62b274e
·
verified ·
1 Parent(s): 06f1a79

Upload Yelp Review Quality Classifier model

Browse files
README.md CHANGED
@@ -1,11 +1,10 @@
1
  ---
2
  library_name: transformers
3
- license: apache-2.0
4
- base_model: distilbert-base-uncased
5
  tags:
 
 
 
6
  - generated_from_trainer
7
- metrics:
8
- - accuracy
9
  model-index:
10
  - name: yelpreviewproject
11
  results: []
@@ -16,10 +15,7 @@ should probably proofread and complete it, then remove this comment. -->
16
 
17
  # yelpreviewproject
18
 
19
- This model is a fine-tuned version of [distilbert-base-uncased](https://huggingface.co/distilbert-base-uncased) on the None dataset.
20
- It achieves the following results on the evaluation set:
21
- - Loss: 0.4755
22
- - Accuracy: 0.764
23
 
24
  ## Model description
25
 
@@ -38,21 +34,13 @@ More information needed
38
  ### Training hyperparameters
39
 
40
  The following hyperparameters were used during training:
41
- - learning_rate: 2e-05
42
- - train_batch_size: 16
43
- - eval_batch_size: 16
44
  - seed: 42
45
  - optimizer: Use OptimizerNames.ADAMW_TORCH_FUSED with betas=(0.9,0.999) and epsilon=1e-08 and optimizer_args=No additional optimizer arguments
46
  - lr_scheduler_type: linear
47
- - num_epochs: 2
48
-
49
- ### Training results
50
-
51
- | Training Loss | Epoch | Step | Validation Loss | Accuracy |
52
- |:-------------:|:-----:|:----:|:---------------:|:--------:|
53
- | 0.554 | 1.0 | 500 | 0.4801 | 0.766 |
54
- | 0.4343 | 2.0 | 1000 | 0.4755 | 0.764 |
55
-
56
 
57
  ### Framework versions
58
 
 
1
  ---
2
  library_name: transformers
 
 
3
  tags:
4
+ - text-classification
5
+ - yelp
6
+ - review-quality
7
  - generated_from_trainer
 
 
8
  model-index:
9
  - name: yelpreviewproject
10
  results: []
 
15
 
16
  # yelpreviewproject
17
 
18
+ This model was trained from scratch on an unknown dataset.
 
 
 
19
 
20
  ## Model description
21
 
 
34
  ### Training hyperparameters
35
 
36
  The following hyperparameters were used during training:
37
+ - learning_rate: 5e-05
38
+ - train_batch_size: 8
39
+ - eval_batch_size: 8
40
  - seed: 42
41
  - optimizer: Use OptimizerNames.ADAMW_TORCH_FUSED with betas=(0.9,0.999) and epsilon=1e-08 and optimizer_args=No additional optimizer arguments
42
  - lr_scheduler_type: linear
43
+ - num_epochs: 3.0
 
 
 
 
 
 
 
 
44
 
45
  ### Framework versions
46
 
config.json CHANGED
@@ -9,15 +9,15 @@
9
  "dtype": "float32",
10
  "hidden_dim": 3072,
11
  "id2label": {
12
- "0": "LABEL_0",
13
- "1": "LABEL_1",
14
- "2": "LABEL_2"
15
  },
16
  "initializer_range": 0.02,
17
  "label2id": {
18
- "LABEL_0": 0,
19
- "LABEL_1": 1,
20
- "LABEL_2": 2
21
  },
22
  "max_position_embeddings": 512,
23
  "model_type": "distilbert",
 
9
  "dtype": "float32",
10
  "hidden_dim": 3072,
11
  "id2label": {
12
+ "0": "Low Quality",
13
+ "1": "Medium Quality",
14
+ "2": "High Quality"
15
  },
16
  "initializer_range": 0.02,
17
  "label2id": {
18
+ "High Quality": 2,
19
+ "Low Quality": 0,
20
+ "Medium Quality": 1
21
  },
22
  "max_position_embeddings": 512,
23
  "model_type": "distilbert",
model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:7b1107f608efed4a27515021da3f85e644c5c8038b7da10efea5ad2a20fccf87
3
  size 267835644
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:fa77a23a3e13990bfe4e59f4a9c32318ee0d5c8c2ca7d926fff3fa45295dc960
3
  size 267835644
special_tokens_map.json CHANGED
@@ -1,7 +1,37 @@
1
  {
2
- "cls_token": "[CLS]",
3
- "mask_token": "[MASK]",
4
- "pad_token": "[PAD]",
5
- "sep_token": "[SEP]",
6
- "unk_token": "[UNK]"
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
7
  }
 
1
  {
2
+ "cls_token": {
3
+ "content": "[CLS]",
4
+ "lstrip": false,
5
+ "normalized": false,
6
+ "rstrip": false,
7
+ "single_word": false
8
+ },
9
+ "mask_token": {
10
+ "content": "[MASK]",
11
+ "lstrip": false,
12
+ "normalized": false,
13
+ "rstrip": false,
14
+ "single_word": false
15
+ },
16
+ "pad_token": {
17
+ "content": "[PAD]",
18
+ "lstrip": false,
19
+ "normalized": false,
20
+ "rstrip": false,
21
+ "single_word": false
22
+ },
23
+ "sep_token": {
24
+ "content": "[SEP]",
25
+ "lstrip": false,
26
+ "normalized": false,
27
+ "rstrip": false,
28
+ "single_word": false
29
+ },
30
+ "unk_token": {
31
+ "content": "[UNK]",
32
+ "lstrip": false,
33
+ "normalized": false,
34
+ "rstrip": false,
35
+ "single_word": false
36
+ }
37
  }
tokenizer_config.json CHANGED
@@ -46,11 +46,15 @@
46
  "do_lower_case": true,
47
  "extra_special_tokens": {},
48
  "mask_token": "[MASK]",
 
49
  "model_max_length": 512,
50
  "pad_token": "[PAD]",
51
  "sep_token": "[SEP]",
 
52
  "strip_accents": null,
53
  "tokenize_chinese_chars": true,
54
  "tokenizer_class": "DistilBertTokenizer",
 
 
55
  "unk_token": "[UNK]"
56
  }
 
46
  "do_lower_case": true,
47
  "extra_special_tokens": {},
48
  "mask_token": "[MASK]",
49
+ "max_length": 256,
50
  "model_max_length": 512,
51
  "pad_token": "[PAD]",
52
  "sep_token": "[SEP]",
53
+ "stride": 0,
54
  "strip_accents": null,
55
  "tokenize_chinese_chars": true,
56
  "tokenizer_class": "DistilBertTokenizer",
57
+ "truncation_side": "right",
58
+ "truncation_strategy": "longest_first",
59
  "unk_token": "[UNK]"
60
  }
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:fd61813d8fff8b6d778a1541df9db289f9388b7aa1683e854839d501e9cf2136
3
- size 5841
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3342e767af86d6c7083470e3225e91b8fa90a9ee63018426d33197c0c808eff6
3
+ size 5905