AmberYifan commited on
Commit
776c30b
·
verified ·
1 Parent(s): a82b679

capmix single-domain code_random_b1000_s0 (marin-8b-base)

Browse files
README.md CHANGED
@@ -16,7 +16,7 @@ should probably proofread and complete it, then remove this comment. -->
16
 
17
  # marin-8b-base_code_random_b1000_s0
18
 
19
- This model is a fine-tuned version of [marin-community/marin-8b-base](https://huggingface.co/marin-community/marin-8b-base) on the capsd_marin-8b-base-n80000-code__mix_code_random_b1000_s0 dataset.
20
 
21
  ## Model description
22
 
 
16
 
17
  # marin-8b-base_code_random_b1000_s0
18
 
19
+ This model is a fine-tuned version of [marin-community/marin-8b-base](https://huggingface.co/marin-community/marin-8b-base) on the capsd_marin-8b-base-n80000-opc__mix_code_random_b1000_s0 dataset.
20
 
21
  ## Model description
22
 
all_results.json CHANGED
@@ -1,8 +1,8 @@
1
  {
2
  "epoch": 1.0,
3
- "total_flos": 591742894080.0,
4
- "train_loss": 0.3349860608577728,
5
- "train_runtime": 513.1196,
6
- "train_samples_per_second": 1.949,
7
- "train_steps_per_second": 0.031
8
  }
 
1
  {
2
  "epoch": 1.0,
3
+ "total_flos": 293928960000.0,
4
+ "train_loss": 0.44835618138313293,
5
+ "train_runtime": 505.4764,
6
+ "train_samples_per_second": 1.978,
7
+ "train_steps_per_second": 0.032
8
  }
checkpoint-16/model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:654864173588567cdccafcc15328cfae9dd21795fa0f4f1ff79c051433f4dc37
3
  size 16060556616
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:890d393c33bb08577646ebfd9e4fccef9a85963af34c7f8008bd615f67df865f
3
  size 16060556616
checkpoint-16/trainer_state.json CHANGED
@@ -26,7 +26,7 @@
26
  "attributes": {}
27
  }
28
  },
29
- "total_flos": 591742894080.0,
30
  "train_batch_size": 2,
31
  "trial_name": null,
32
  "trial_params": null
 
26
  "attributes": {}
27
  }
28
  },
29
+ "total_flos": 293928960000.0,
30
  "train_batch_size": 2,
31
  "trial_name": null,
32
  "trial_params": null
checkpoint-16/training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:9e2f90276c855ffd23f7dc7ce13101ddaf5c4523206c9baeb55816b8930bd272
3
  size 7761
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:144ab5073ac7c48be0a50e41fdda757cc8d9ffd5204d108bac9f585cf5726e43
3
  size 7761
model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:654864173588567cdccafcc15328cfae9dd21795fa0f4f1ff79c051433f4dc37
3
  size 16060556616
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:890d393c33bb08577646ebfd9e4fccef9a85963af34c7f8008bd615f67df865f
3
  size 16060556616
train_results.json CHANGED
@@ -1,8 +1,8 @@
1
  {
2
  "epoch": 1.0,
3
- "total_flos": 591742894080.0,
4
- "train_loss": 0.3349860608577728,
5
- "train_runtime": 513.1196,
6
- "train_samples_per_second": 1.949,
7
- "train_steps_per_second": 0.031
8
  }
 
1
  {
2
  "epoch": 1.0,
3
+ "total_flos": 293928960000.0,
4
+ "train_loss": 0.44835618138313293,
5
+ "train_runtime": 505.4764,
6
+ "train_samples_per_second": 1.978,
7
+ "train_steps_per_second": 0.032
8
  }
trainer_log.jsonl CHANGED
@@ -1 +1 @@
1
- {"current_steps": 16, "total_steps": 16, "epoch": 1.0, "percentage": 100.0, "elapsed_time": "0:08:33", "remaining_time": "0:00:00"}
 
1
+ {"current_steps": 16, "total_steps": 16, "epoch": 1.0, "percentage": 100.0, "elapsed_time": "0:08:25", "remaining_time": "0:00:00"}
trainer_state.json CHANGED
@@ -12,11 +12,11 @@
12
  {
13
  "epoch": 1.0,
14
  "step": 16,
15
- "total_flos": 591742894080.0,
16
- "train_loss": 0.3349860608577728,
17
- "train_runtime": 513.1196,
18
- "train_samples_per_second": 1.949,
19
- "train_steps_per_second": 0.031
20
  }
21
  ],
22
  "logging_steps": 50,
@@ -36,7 +36,7 @@
36
  "attributes": {}
37
  }
38
  },
39
- "total_flos": 591742894080.0,
40
  "train_batch_size": 2,
41
  "trial_name": null,
42
  "trial_params": null
 
12
  {
13
  "epoch": 1.0,
14
  "step": 16,
15
+ "total_flos": 293928960000.0,
16
+ "train_loss": 0.44835618138313293,
17
+ "train_runtime": 505.4764,
18
+ "train_samples_per_second": 1.978,
19
+ "train_steps_per_second": 0.032
20
  }
21
  ],
22
  "logging_steps": 50,
 
36
  "attributes": {}
37
  }
38
  },
39
+ "total_flos": 293928960000.0,
40
  "train_batch_size": 2,
41
  "trial_name": null,
42
  "trial_params": null
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:9e2f90276c855ffd23f7dc7ce13101ddaf5c4523206c9baeb55816b8930bd272
3
  size 7761
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:144ab5073ac7c48be0a50e41fdda757cc8d9ffd5204d108bac9f585cf5726e43
3
  size 7761