Training in progress, step 10

Files changed (6) hide show

adapter_config.json CHANGED Viewed

@@ -3,7 +3,7 @@
   "alpha_pattern": {},
   "arrow_config": null,
   "auto_mapping": null,
-  "base_model_name_or_path": "meta-llama/Llama-3.1-8B-Instruct",
   "bias": "none",
   "corda_config": null,
   "ensure_weight_tying": false,
@@ -29,13 +29,13 @@
   "rank_pattern": {},
   "revision": null,
   "target_modules": [
-    "down_proj",
-    "gate_proj",
     "v_proj",
     "k_proj",
     "up_proj",
-    "o_proj",
-    "q_proj"
   ],
   "target_parameters": null,
   "task_type": "CAUSAL_LM",

   "alpha_pattern": {},
   "arrow_config": null,
   "auto_mapping": null,
+  "base_model_name_or_path": "meta-llama/Llama-3.2-3B",
   "bias": "none",
   "corda_config": null,
   "ensure_weight_tying": false,
   "rank_pattern": {},
   "revision": null,
   "target_modules": [
     "v_proj",
+    "q_proj",
+    "o_proj",
     "k_proj",
     "up_proj",
+    "down_proj",
+    "gate_proj"
   ],
   "target_parameters": null,
   "task_type": "CAUSAL_LM",

adapter_model.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:7dbd4d8cb3854070ae41b2cced025798646677468d5841b80a19d20fc089887b
-size 335604696

 version https://git-lfs.github.com/spec/v1
+oid sha256:6760d31fbf3d5a6343036b485d27a3a76c64b35026af34d485c514d2435acb38
+size 194563400

runs/Dec17_09-52-46_c45610aa664c/events.out.tfevents.1765965171.c45610aa664c.1078.0 ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:d4a2aebad239b7de0ff9579e6aecca463feb26704fc70425cc0422166e785ab0
+size 6251

special_tokens_map.json CHANGED Viewed

@@ -7,11 +7,11 @@
     "single_word": false
   },
   "eos_token": {
-    "content": "<|eot_id|>",
     "lstrip": false,
     "normalized": false,
     "rstrip": false,
     "single_word": false
   },
-  "pad_token": "<|eot_id|>"
 }

     "single_word": false
   },
   "eos_token": {
+    "content": "<|end_of_text|>",
     "lstrip": false,
     "normalized": false,
     "rstrip": false,
     "single_word": false
   },
+  "pad_token": "<|end_of_text|>"
 }

tokenizer_config.json CHANGED Viewed

@@ -2051,13 +2051,13 @@
   },
   "bos_token": "<|begin_of_text|>",
   "clean_up_tokenization_spaces": true,
-  "eos_token": "<|eot_id|>",
   "extra_special_tokens": {},
   "model_input_names": [
     "input_ids",
     "attention_mask"
   ],
   "model_max_length": 131072,
-  "pad_token": "<|eot_id|>",
   "tokenizer_class": "PreTrainedTokenizerFast"
 }

   },
   "bos_token": "<|begin_of_text|>",
   "clean_up_tokenization_spaces": true,
+  "eos_token": "<|end_of_text|>",
   "extra_special_tokens": {},
   "model_input_names": [
     "input_ids",
     "attention_mask"
   ],
   "model_max_length": 131072,
+  "pad_token": "<|end_of_text|>",
   "tokenizer_class": "PreTrainedTokenizerFast"
 }

training_args.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:33f7833fc984ed42891ed683c63672751de4e1502946c1ab8ea9df9d21885712
 size 6097

 version https://git-lfs.github.com/spec/v1
+oid sha256:3b9ff7fd9432f2963474cfa9bf43b7ada94fb55ffb84095a19d5f167ecd5fd0a
 size 6097