YagiASAFAS commited on Jun 14, 2025

Commit

d678842

verified ·

1 Parent(s): de7723e

Training in progress, epoch 1

Browse files

Files changed (34) hide show

model.safetensors +1 -1
run-2/checkpoint-1000/config.json +93 -0
run-2/checkpoint-1000/model.safetensors +3 -0
run-2/checkpoint-1000/optimizer.pt +3 -0
run-2/checkpoint-1000/rng_state.pth +3 -0
run-2/checkpoint-1000/scaler.pt +3 -0
run-2/checkpoint-1000/scheduler.pt +3 -0
run-2/checkpoint-1000/trainer_state.json +221 -0
run-2/checkpoint-1000/training_args.bin +3 -0
run-2/checkpoint-1250/config.json +93 -0
run-2/checkpoint-1250/model.safetensors +3 -0
run-2/checkpoint-1250/optimizer.pt +3 -0
run-2/checkpoint-1250/rng_state.pth +3 -0
run-2/checkpoint-1250/scaler.pt +3 -0
run-2/checkpoint-1250/scheduler.pt +3 -0
run-2/checkpoint-1250/trainer_state.json +261 -0
run-2/checkpoint-1250/training_args.bin +3 -0
run-2/checkpoint-750/config.json +93 -0
run-2/checkpoint-750/model.safetensors +3 -0
run-2/checkpoint-750/optimizer.pt +3 -0
run-2/checkpoint-750/rng_state.pth +3 -0
run-2/checkpoint-750/scaler.pt +3 -0
run-2/checkpoint-750/scheduler.pt +3 -0
run-2/checkpoint-750/trainer_state.json +174 -0
run-2/checkpoint-750/training_args.bin +3 -0
run-3/checkpoint-125/config.json +93 -0
run-3/checkpoint-125/model.safetensors +3 -0
run-3/checkpoint-125/optimizer.pt +3 -0
run-3/checkpoint-125/rng_state.pth +3 -0
run-3/checkpoint-125/scaler.pt +3 -0
run-3/checkpoint-125/scheduler.pt +3 -0
run-3/checkpoint-125/trainer_state.json +80 -0
run-3/checkpoint-125/training_args.bin +3 -0
training_args.bin +1 -1

model.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:58fbf83a5dfd0f1b72fd9a4dffb486738a2a0c7aeb0a909d8c62224393145608
 size 438050928

 version https://git-lfs.github.com/spec/v1
+oid sha256:a1eb23c4286bb70d60d37d169b4eb4cdef72b61f202122c30bc3dc70490c88d5
 size 438050928

run-2/checkpoint-1000/config.json ADDED Viewed

	@@ -0,0 +1,93 @@

+{
+  "architectures": [
+    "BertForSequenceClassification"
+  ],
+  "attention_probs_dropout_prob": 0.1,
+  "classifier_dropout": null,
+  "gradient_checkpointing": false,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
+  "hidden_size": 768,
+  "id2label": {
+    "0": "LABEL_0",
+    "1": "LABEL_1",
+    "2": "LABEL_2",
+    "3": "LABEL_3",
+    "4": "LABEL_4",
+    "5": "LABEL_5",
+    "6": "LABEL_6",
+    "7": "LABEL_7",
+    "8": "LABEL_8",
+    "9": "LABEL_9",
+    "10": "LABEL_10",
+    "11": "LABEL_11",
+    "12": "LABEL_12",
+    "13": "LABEL_13",
+    "14": "LABEL_14",
+    "15": "LABEL_15",
+    "16": "LABEL_16",
+    "17": "LABEL_17",
+    "18": "LABEL_18",
+    "19": "LABEL_19",
+    "20": "LABEL_20",
+    "21": "LABEL_21",
+    "22": "LABEL_22",
+    "23": "LABEL_23",
+    "24": "LABEL_24",
+    "25": "LABEL_25",
+    "26": "LABEL_26",
+    "27": "LABEL_27",
+    "28": "LABEL_28",
+    "29": "LABEL_29",
+    "30": "LABEL_30",
+    "31": "LABEL_31"
+  },
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "label2id": {
+    "LABEL_0": 0,
+    "LABEL_1": 1,
+    "LABEL_10": 10,
+    "LABEL_11": 11,
+    "LABEL_12": 12,
+    "LABEL_13": 13,
+    "LABEL_14": 14,
+    "LABEL_15": 15,
+    "LABEL_16": 16,
+    "LABEL_17": 17,
+    "LABEL_18": 18,
+    "LABEL_19": 19,
+    "LABEL_2": 2,
+    "LABEL_20": 20,
+    "LABEL_21": 21,
+    "LABEL_22": 22,
+    "LABEL_23": 23,
+    "LABEL_24": 24,
+    "LABEL_25": 25,
+    "LABEL_26": 26,
+    "LABEL_27": 27,
+    "LABEL_28": 28,
+    "LABEL_29": 29,
+    "LABEL_3": 3,
+    "LABEL_30": 30,
+    "LABEL_31": 31,
+    "LABEL_4": 4,
+    "LABEL_5": 5,
+    "LABEL_6": 6,
+    "LABEL_7": 7,
+    "LABEL_8": 8,
+    "LABEL_9": 9
+  },
+  "layer_norm_eps": 1e-12,
+  "max_position_embeddings": 512,
+  "model_type": "bert",
+  "num_attention_heads": 12,
+  "num_hidden_layers": 12,
+  "pad_token_id": 0,
+  "position_embedding_type": "absolute",
+  "torch_dtype": "float32",
+  "transformers_version": "4.52.4",
+  "type_vocab_size": 2,
+  "use_cache": true,
+  "vocab_size": 30522
+}

run-2/checkpoint-1000/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:3cd644eb1cba868925e4e9d98f38fc0bbcbf03d4236fd6168287da38fac0515e
+size 438050928

run-2/checkpoint-1000/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d1a434390b13116d75a2374e8d05c4512aa2d9bde83fb1984269025dc57922ec
+size 876222970

run-2/checkpoint-1000/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:895ec29675e68d1b3c1927e42320e3a12af8a184e10c106999c38ddae836cafc
+size 14244

run-2/checkpoint-1000/scaler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:9d8fdcd0311eba9854fff738038ed4c1a269832665b4d88ba4e4e3d02a1a7e0e
+size 988

run-2/checkpoint-1000/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:973ca42b92d2145c5e30a6c46bf31f069380c8349a7b567f0da38d34836145f3
+size 1064

run-2/checkpoint-1000/trainer_state.json ADDED Viewed

	@@ -0,0 +1,221 @@

+{
+  "best_global_step": 1000,
+  "best_metric": 0.9414763117147434,
+  "best_model_checkpoint": "./results/run-2/checkpoint-1000",
+  "epoch": 4.0,
+  "eval_steps": 500,
+  "global_step": 1000,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 0.4,
+      "grad_norm": 1.9651439189910889,
+      "learning_rate": 1.386e-05,
+      "loss": 1.62,
+      "step": 100
+    },
+    {
+      "epoch": 0.8,
+      "grad_norm": 1.7124375104904175,
+      "learning_rate": 2.7859999999999998e-05,
+      "loss": 0.6715,
+      "step": 200
+    },
+    {
+      "epoch": 1.0,
+      "eval_clientelism_accuracy": 0.906,
+      "eval_clientelism_f1": 0.897770957127279,
+      "eval_discipline_among_poor_accuracy": 0.968,
+      "eval_discipline_among_poor_f1": 0.9615640823822642,
+      "eval_economic_policy_accuracy": 0.875,
+      "eval_economic_policy_f1": 0.860741831908876,
+      "eval_loss": 0.47158536314964294,
+      "eval_marcos_duterte_alliance_accuracy": 0.846,
+      "eval_marcos_duterte_alliance_f1": 0.839269429776376,
+      "eval_overall_accuracy": 0.8419375,
+      "eval_overall_f1": 0.8147702755950403,
+      "eval_populism_accuracy": 0.645,
+      "eval_populism_f1": 0.5639441230465639,
+      "eval_regionalism_accuracy": 0.963,
+      "eval_regionalism_f1": 0.9482946255564284,
+      "eval_runtime": 2.5786,
+      "eval_samples_per_second": 775.6,
+      "eval_security_accuracy": 0.8275,
+      "eval_security_f1": 0.8006718398296887,
+      "eval_steps_per_second": 48.475,
+      "eval_uniteam_positive_campaign_accuracy": 0.705,
+      "eval_uniteam_positive_campaign_f1": 0.645905315132846,
+      "step": 250
+    },
+    {
+      "epoch": 1.2,
+      "grad_norm": 1.057538390159607,
+      "learning_rate": 4.1859999999999996e-05,
+      "loss": 0.4761,
+      "step": 300
+    },
+    {
+      "epoch": 1.6,
+      "grad_norm": 1.3597612380981445,
+      "learning_rate": 5.586e-05,
+      "loss": 0.3424,
+      "step": 400
+    },
+    {
+      "epoch": 2.0,
+      "grad_norm": 1.4431722164154053,
+      "learning_rate": 6.986e-05,
+      "loss": 0.2992,
+      "step": 500
+    },
+    {
+      "epoch": 2.0,
+      "eval_clientelism_accuracy": 0.949,
+      "eval_clientelism_f1": 0.9452056048035566,
+      "eval_discipline_among_poor_accuracy": 0.975,
+      "eval_discipline_among_poor_f1": 0.9713251858526694,
+      "eval_economic_policy_accuracy": 0.946,
+      "eval_economic_policy_f1": 0.9447796444638326,
+      "eval_loss": 0.2847972512245178,
+      "eval_marcos_duterte_alliance_accuracy": 0.905,
+      "eval_marcos_duterte_alliance_f1": 0.8931059723613894,
+      "eval_overall_accuracy": 0.9166874999999999,
+      "eval_overall_f1": 0.9105488176533825,
+      "eval_populism_accuracy": 0.7845,
+      "eval_populism_f1": 0.7756216085131521,
+      "eval_regionalism_accuracy": 0.977,
+      "eval_regionalism_f1": 0.9741474164644311,
+      "eval_runtime": 2.6526,
+      "eval_samples_per_second": 753.967,
+      "eval_security_accuracy": 0.943,
+      "eval_security_f1": 0.934999506345517,
+      "eval_steps_per_second": 47.123,
+      "eval_uniteam_positive_campaign_accuracy": 0.854,
+      "eval_uniteam_positive_campaign_f1": 0.8452056024225119,
+      "step": 500
+    },
+    {
+      "epoch": 2.4,
+      "grad_norm": 1.2808622121810913,
+      "learning_rate": 6.0759999999999994e-05,
+      "loss": 0.2152,
+      "step": 600
+    },
+    {
+      "epoch": 2.8,
+      "grad_norm": 0.9503875970840454,
+      "learning_rate": 5.142666666666667e-05,
+      "loss": 0.1943,
+      "step": 700
+    },
+    {
+      "epoch": 3.0,
+      "eval_clientelism_accuracy": 0.954,
+      "eval_clientelism_f1": 0.9527297868951113,
+      "eval_discipline_among_poor_accuracy": 0.977,
+      "eval_discipline_among_poor_f1": 0.975000371636688,
+      "eval_economic_policy_accuracy": 0.9495,
+      "eval_economic_policy_f1": 0.9493495071116878,
+      "eval_loss": 0.2306518405675888,
+      "eval_marcos_duterte_alliance_accuracy": 0.9325,
+      "eval_marcos_duterte_alliance_f1": 0.9269038834173264,
+      "eval_overall_accuracy": 0.9380625,
+      "eval_overall_f1": 0.9368632136576173,
+      "eval_populism_accuracy": 0.8745,
+      "eval_populism_f1": 0.8748340797000174,
+      "eval_regionalism_accuracy": 0.9735,
+      "eval_regionalism_f1": 0.9749884932130427,
+      "eval_runtime": 2.7943,
+      "eval_samples_per_second": 715.749,
+      "eval_security_accuracy": 0.953,
+      "eval_security_f1": 0.9496585291837848,
+      "eval_steps_per_second": 44.734,
+      "eval_uniteam_positive_campaign_accuracy": 0.8905,
+      "eval_uniteam_positive_campaign_f1": 0.8914410581032797,
+      "step": 750
+    },
+    {
+      "epoch": 3.2,
+      "grad_norm": 0.7700262069702148,
+      "learning_rate": 4.209333333333333e-05,
+      "loss": 0.155,
+      "step": 800
+    },
+    {
+      "epoch": 3.6,
+      "grad_norm": 0.646346390247345,
+      "learning_rate": 3.276e-05,
+      "loss": 0.1326,
+      "step": 900
+    },
+    {
+      "epoch": 4.0,
+      "grad_norm": 0.8825557231903076,
+      "learning_rate": 2.3426666666666664e-05,
+      "loss": 0.1271,
+      "step": 1000
+    },
+    {
+      "epoch": 4.0,
+      "eval_clientelism_accuracy": 0.953,
+      "eval_clientelism_f1": 0.9530642870880963,
+      "eval_discipline_among_poor_accuracy": 0.9735,
+      "eval_discipline_among_poor_f1": 0.9718915263402587,
+      "eval_economic_policy_accuracy": 0.9575,
+      "eval_economic_policy_f1": 0.9571061555033147,
+      "eval_loss": 0.21463848650455475,
+      "eval_marcos_duterte_alliance_accuracy": 0.935,
+      "eval_marcos_duterte_alliance_f1": 0.9294683134335573,
+      "eval_overall_accuracy": 0.9431875,
+      "eval_overall_f1": 0.9414763117147434,
+      "eval_populism_accuracy": 0.889,
+      "eval_populism_f1": 0.8882672567299156,
+      "eval_regionalism_accuracy": 0.979,
+      "eval_regionalism_f1": 0.9783012417980759,
+      "eval_runtime": 2.5606,
+      "eval_samples_per_second": 781.072,
+      "eval_security_accuracy": 0.9555,
+      "eval_security_f1": 0.9523932010152734,
+      "eval_steps_per_second": 48.817,
+      "eval_uniteam_positive_campaign_accuracy": 0.903,
+      "eval_uniteam_positive_campaign_f1": 0.9013185118094562,
+      "step": 1000
+    }
+  ],
+  "logging_steps": 100,
+  "max_steps": 1250,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 5,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "EarlyStoppingCallback": {
+      "args": {
+        "early_stopping_patience": 2,
+        "early_stopping_threshold": 0.0
+      },
+      "attributes": {
+        "early_stopping_patience_counter": 0
+      }
+    },
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": false
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 8421821644800000.0,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "gradient_accumulation_steps": 2,
+    "learning_rate": 7e-05,
+    "num_train_epochs": 5
+  }
+}

run-2/checkpoint-1000/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:00db87b530f828066f2ef5374605ff7c067cf8bcaf6d07e1ad2162b4b005b6c2
+size 5368

run-2/checkpoint-1250/config.json ADDED Viewed

	@@ -0,0 +1,93 @@

+{
+  "architectures": [
+    "BertForSequenceClassification"
+  ],
+  "attention_probs_dropout_prob": 0.1,
+  "classifier_dropout": null,
+  "gradient_checkpointing": false,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
+  "hidden_size": 768,
+  "id2label": {
+    "0": "LABEL_0",
+    "1": "LABEL_1",
+    "2": "LABEL_2",
+    "3": "LABEL_3",
+    "4": "LABEL_4",
+    "5": "LABEL_5",
+    "6": "LABEL_6",
+    "7": "LABEL_7",
+    "8": "LABEL_8",
+    "9": "LABEL_9",
+    "10": "LABEL_10",
+    "11": "LABEL_11",
+    "12": "LABEL_12",
+    "13": "LABEL_13",
+    "14": "LABEL_14",
+    "15": "LABEL_15",
+    "16": "LABEL_16",
+    "17": "LABEL_17",
+    "18": "LABEL_18",
+    "19": "LABEL_19",
+    "20": "LABEL_20",
+    "21": "LABEL_21",
+    "22": "LABEL_22",
+    "23": "LABEL_23",
+    "24": "LABEL_24",
+    "25": "LABEL_25",
+    "26": "LABEL_26",
+    "27": "LABEL_27",
+    "28": "LABEL_28",
+    "29": "LABEL_29",
+    "30": "LABEL_30",
+    "31": "LABEL_31"
+  },
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "label2id": {
+    "LABEL_0": 0,
+    "LABEL_1": 1,
+    "LABEL_10": 10,
+    "LABEL_11": 11,
+    "LABEL_12": 12,
+    "LABEL_13": 13,
+    "LABEL_14": 14,
+    "LABEL_15": 15,
+    "LABEL_16": 16,
+    "LABEL_17": 17,
+    "LABEL_18": 18,
+    "LABEL_19": 19,
+    "LABEL_2": 2,
+    "LABEL_20": 20,
+    "LABEL_21": 21,
+    "LABEL_22": 22,
+    "LABEL_23": 23,
+    "LABEL_24": 24,
+    "LABEL_25": 25,
+    "LABEL_26": 26,
+    "LABEL_27": 27,
+    "LABEL_28": 28,
+    "LABEL_29": 29,
+    "LABEL_3": 3,
+    "LABEL_30": 30,
+    "LABEL_31": 31,
+    "LABEL_4": 4,
+    "LABEL_5": 5,
+    "LABEL_6": 6,
+    "LABEL_7": 7,
+    "LABEL_8": 8,
+    "LABEL_9": 9
+  },
+  "layer_norm_eps": 1e-12,
+  "max_position_embeddings": 512,
+  "model_type": "bert",
+  "num_attention_heads": 12,
+  "num_hidden_layers": 12,
+  "pad_token_id": 0,
+  "position_embedding_type": "absolute",
+  "torch_dtype": "float32",
+  "transformers_version": "4.52.4",
+  "type_vocab_size": 2,
+  "use_cache": true,
+  "vocab_size": 30522
+}

run-2/checkpoint-1250/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6e2969858d95c2dce93a840aa08aa3087fec871ad83392c711ccd5defb3d9cb5
+size 438050928

run-2/checkpoint-1250/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:c743a2d689ece1d562a42b7f603e00ae333a1ffbab25d1ec8df1d26b84e872a8
+size 876222970

run-2/checkpoint-1250/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e075e1532db1c17ec921b935a69934372058cf583818c89d05c8bc844c362dc2
+size 14244

run-2/checkpoint-1250/scaler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:01f1986269c74ce38fd01c56e0770c176375c55daeed6953aaf84f06cae775ba
+size 988

run-2/checkpoint-1250/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:be8f1f74b2ddc5f3bcefc3afc5761b8b1e0d3a04d708169af990768b4a718d20
+size 1064

run-2/checkpoint-1250/trainer_state.json ADDED Viewed

	@@ -0,0 +1,261 @@

+{
+  "best_global_step": 1250,
+  "best_metric": 0.9440323611457324,
+  "best_model_checkpoint": "./results/run-2/checkpoint-1250",
+  "epoch": 5.0,
+  "eval_steps": 500,
+  "global_step": 1250,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 0.4,
+      "grad_norm": 1.9651439189910889,
+      "learning_rate": 1.386e-05,
+      "loss": 1.62,
+      "step": 100
+    },
+    {
+      "epoch": 0.8,
+      "grad_norm": 1.7124375104904175,
+      "learning_rate": 2.7859999999999998e-05,
+      "loss": 0.6715,
+      "step": 200
+    },
+    {
+      "epoch": 1.0,
+      "eval_clientelism_accuracy": 0.906,
+      "eval_clientelism_f1": 0.897770957127279,
+      "eval_discipline_among_poor_accuracy": 0.968,
+      "eval_discipline_among_poor_f1": 0.9615640823822642,
+      "eval_economic_policy_accuracy": 0.875,
+      "eval_economic_policy_f1": 0.860741831908876,
+      "eval_loss": 0.47158536314964294,
+      "eval_marcos_duterte_alliance_accuracy": 0.846,
+      "eval_marcos_duterte_alliance_f1": 0.839269429776376,
+      "eval_overall_accuracy": 0.8419375,
+      "eval_overall_f1": 0.8147702755950403,
+      "eval_populism_accuracy": 0.645,
+      "eval_populism_f1": 0.5639441230465639,
+      "eval_regionalism_accuracy": 0.963,
+      "eval_regionalism_f1": 0.9482946255564284,
+      "eval_runtime": 2.5786,
+      "eval_samples_per_second": 775.6,
+      "eval_security_accuracy": 0.8275,
+      "eval_security_f1": 0.8006718398296887,
+      "eval_steps_per_second": 48.475,
+      "eval_uniteam_positive_campaign_accuracy": 0.705,
+      "eval_uniteam_positive_campaign_f1": 0.645905315132846,
+      "step": 250
+    },
+    {
+      "epoch": 1.2,
+      "grad_norm": 1.057538390159607,
+      "learning_rate": 4.1859999999999996e-05,
+      "loss": 0.4761,
+      "step": 300
+    },
+    {
+      "epoch": 1.6,
+      "grad_norm": 1.3597612380981445,
+      "learning_rate": 5.586e-05,
+      "loss": 0.3424,
+      "step": 400
+    },
+    {
+      "epoch": 2.0,
+      "grad_norm": 1.4431722164154053,
+      "learning_rate": 6.986e-05,
+      "loss": 0.2992,
+      "step": 500
+    },
+    {
+      "epoch": 2.0,
+      "eval_clientelism_accuracy": 0.949,
+      "eval_clientelism_f1": 0.9452056048035566,
+      "eval_discipline_among_poor_accuracy": 0.975,
+      "eval_discipline_among_poor_f1": 0.9713251858526694,
+      "eval_economic_policy_accuracy": 0.946,
+      "eval_economic_policy_f1": 0.9447796444638326,
+      "eval_loss": 0.2847972512245178,
+      "eval_marcos_duterte_alliance_accuracy": 0.905,
+      "eval_marcos_duterte_alliance_f1": 0.8931059723613894,
+      "eval_overall_accuracy": 0.9166874999999999,
+      "eval_overall_f1": 0.9105488176533825,
+      "eval_populism_accuracy": 0.7845,
+      "eval_populism_f1": 0.7756216085131521,
+      "eval_regionalism_accuracy": 0.977,
+      "eval_regionalism_f1": 0.9741474164644311,
+      "eval_runtime": 2.6526,
+      "eval_samples_per_second": 753.967,
+      "eval_security_accuracy": 0.943,
+      "eval_security_f1": 0.934999506345517,
+      "eval_steps_per_second": 47.123,
+      "eval_uniteam_positive_campaign_accuracy": 0.854,
+      "eval_uniteam_positive_campaign_f1": 0.8452056024225119,
+      "step": 500
+    },
+    {
+      "epoch": 2.4,
+      "grad_norm": 1.2808622121810913,
+      "learning_rate": 6.0759999999999994e-05,
+      "loss": 0.2152,
+      "step": 600
+    },
+    {
+      "epoch": 2.8,
+      "grad_norm": 0.9503875970840454,
+      "learning_rate": 5.142666666666667e-05,
+      "loss": 0.1943,
+      "step": 700
+    },
+    {
+      "epoch": 3.0,
+      "eval_clientelism_accuracy": 0.954,
+      "eval_clientelism_f1": 0.9527297868951113,
+      "eval_discipline_among_poor_accuracy": 0.977,
+      "eval_discipline_among_poor_f1": 0.975000371636688,
+      "eval_economic_policy_accuracy": 0.9495,
+      "eval_economic_policy_f1": 0.9493495071116878,
+      "eval_loss": 0.2306518405675888,
+      "eval_marcos_duterte_alliance_accuracy": 0.9325,
+      "eval_marcos_duterte_alliance_f1": 0.9269038834173264,
+      "eval_overall_accuracy": 0.9380625,
+      "eval_overall_f1": 0.9368632136576173,
+      "eval_populism_accuracy": 0.8745,
+      "eval_populism_f1": 0.8748340797000174,
+      "eval_regionalism_accuracy": 0.9735,
+      "eval_regionalism_f1": 0.9749884932130427,
+      "eval_runtime": 2.7943,
+      "eval_samples_per_second": 715.749,
+      "eval_security_accuracy": 0.953,
+      "eval_security_f1": 0.9496585291837848,
+      "eval_steps_per_second": 44.734,
+      "eval_uniteam_positive_campaign_accuracy": 0.8905,
+      "eval_uniteam_positive_campaign_f1": 0.8914410581032797,
+      "step": 750
+    },
+    {
+      "epoch": 3.2,
+      "grad_norm": 0.7700262069702148,
+      "learning_rate": 4.209333333333333e-05,
+      "loss": 0.155,
+      "step": 800
+    },
+    {
+      "epoch": 3.6,
+      "grad_norm": 0.646346390247345,
+      "learning_rate": 3.276e-05,
+      "loss": 0.1326,
+      "step": 900
+    },
+    {
+      "epoch": 4.0,
+      "grad_norm": 0.8825557231903076,
+      "learning_rate": 2.3426666666666664e-05,
+      "loss": 0.1271,
+      "step": 1000
+    },
+    {
+      "epoch": 4.0,
+      "eval_clientelism_accuracy": 0.953,
+      "eval_clientelism_f1": 0.9530642870880963,
+      "eval_discipline_among_poor_accuracy": 0.9735,
+      "eval_discipline_among_poor_f1": 0.9718915263402587,
+      "eval_economic_policy_accuracy": 0.9575,
+      "eval_economic_policy_f1": 0.9571061555033147,
+      "eval_loss": 0.21463848650455475,
+      "eval_marcos_duterte_alliance_accuracy": 0.935,
+      "eval_marcos_duterte_alliance_f1": 0.9294683134335573,
+      "eval_overall_accuracy": 0.9431875,
+      "eval_overall_f1": 0.9414763117147434,
+      "eval_populism_accuracy": 0.889,
+      "eval_populism_f1": 0.8882672567299156,
+      "eval_regionalism_accuracy": 0.979,
+      "eval_regionalism_f1": 0.9783012417980759,
+      "eval_runtime": 2.5606,
+      "eval_samples_per_second": 781.072,
+      "eval_security_accuracy": 0.9555,
+      "eval_security_f1": 0.9523932010152734,
+      "eval_steps_per_second": 48.817,
+      "eval_uniteam_positive_campaign_accuracy": 0.903,
+      "eval_uniteam_positive_campaign_f1": 0.9013185118094562,
+      "step": 1000
+    },
+    {
+      "epoch": 4.4,
+      "grad_norm": 0.6359074711799622,
+      "learning_rate": 1.4093333333333333e-05,
+      "loss": 0.1052,
+      "step": 1100
+    },
+    {
+      "epoch": 4.8,
+      "grad_norm": 0.7204127907752991,
+      "learning_rate": 4.76e-06,
+      "loss": 0.093,
+      "step": 1200
+    },
+    {
+      "epoch": 5.0,
+      "eval_clientelism_accuracy": 0.9605,
+      "eval_clientelism_f1": 0.9598378695080977,
+      "eval_discipline_among_poor_accuracy": 0.9775,
+      "eval_discipline_among_poor_f1": 0.9752991563048182,
+      "eval_economic_policy_accuracy": 0.958,
+      "eval_economic_policy_f1": 0.9578000392532403,
+      "eval_loss": 0.20537178218364716,
+      "eval_marcos_duterte_alliance_accuracy": 0.944,
+      "eval_marcos_duterte_alliance_f1": 0.938033209852473,
+      "eval_overall_accuracy": 0.9458125,
+      "eval_overall_f1": 0.9440323611457324,
+      "eval_populism_accuracy": 0.8905,
+      "eval_populism_f1": 0.8894935397931775,
+      "eval_regionalism_accuracy": 0.9815,
+      "eval_regionalism_f1": 0.9812972117328942,
+      "eval_runtime": 2.3936,
+      "eval_samples_per_second": 835.557,
+      "eval_security_accuracy": 0.9545,
+      "eval_security_f1": 0.9512383521837636,
+      "eval_steps_per_second": 52.222,
+      "eval_uniteam_positive_campaign_accuracy": 0.9,
+      "eval_uniteam_positive_campaign_f1": 0.899259510537395,
+      "step": 1250
+    }
+  ],
+  "logging_steps": 100,
+  "max_steps": 1250,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 5,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "EarlyStoppingCallback": {
+      "args": {
+        "early_stopping_patience": 2,
+        "early_stopping_threshold": 0.0
+      },
+      "attributes": {
+        "early_stopping_patience_counter": 0
+      }
+    },
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": true
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 1.010618597376e+16,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "gradient_accumulation_steps": 2,
+    "learning_rate": 7e-05,
+    "num_train_epochs": 5
+  }
+}

run-2/checkpoint-1250/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:00db87b530f828066f2ef5374605ff7c067cf8bcaf6d07e1ad2162b4b005b6c2
+size 5368

run-2/checkpoint-750/config.json ADDED Viewed

	@@ -0,0 +1,93 @@

+{
+  "architectures": [
+    "BertForSequenceClassification"
+  ],
+  "attention_probs_dropout_prob": 0.1,
+  "classifier_dropout": null,
+  "gradient_checkpointing": false,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
+  "hidden_size": 768,
+  "id2label": {
+    "0": "LABEL_0",
+    "1": "LABEL_1",
+    "2": "LABEL_2",
+    "3": "LABEL_3",
+    "4": "LABEL_4",
+    "5": "LABEL_5",
+    "6": "LABEL_6",
+    "7": "LABEL_7",
+    "8": "LABEL_8",
+    "9": "LABEL_9",
+    "10": "LABEL_10",
+    "11": "LABEL_11",
+    "12": "LABEL_12",
+    "13": "LABEL_13",
+    "14": "LABEL_14",
+    "15": "LABEL_15",
+    "16": "LABEL_16",
+    "17": "LABEL_17",
+    "18": "LABEL_18",
+    "19": "LABEL_19",
+    "20": "LABEL_20",
+    "21": "LABEL_21",
+    "22": "LABEL_22",
+    "23": "LABEL_23",
+    "24": "LABEL_24",
+    "25": "LABEL_25",
+    "26": "LABEL_26",
+    "27": "LABEL_27",
+    "28": "LABEL_28",
+    "29": "LABEL_29",
+    "30": "LABEL_30",
+    "31": "LABEL_31"
+  },
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "label2id": {
+    "LABEL_0": 0,
+    "LABEL_1": 1,
+    "LABEL_10": 10,
+    "LABEL_11": 11,
+    "LABEL_12": 12,
+    "LABEL_13": 13,
+    "LABEL_14": 14,
+    "LABEL_15": 15,
+    "LABEL_16": 16,
+    "LABEL_17": 17,
+    "LABEL_18": 18,
+    "LABEL_19": 19,
+    "LABEL_2": 2,
+    "LABEL_20": 20,
+    "LABEL_21": 21,
+    "LABEL_22": 22,
+    "LABEL_23": 23,
+    "LABEL_24": 24,
+    "LABEL_25": 25,
+    "LABEL_26": 26,
+    "LABEL_27": 27,
+    "LABEL_28": 28,
+    "LABEL_29": 29,
+    "LABEL_3": 3,
+    "LABEL_30": 30,
+    "LABEL_31": 31,
+    "LABEL_4": 4,
+    "LABEL_5": 5,
+    "LABEL_6": 6,
+    "LABEL_7": 7,
+    "LABEL_8": 8,
+    "LABEL_9": 9
+  },
+  "layer_norm_eps": 1e-12,
+  "max_position_embeddings": 512,
+  "model_type": "bert",
+  "num_attention_heads": 12,
+  "num_hidden_layers": 12,
+  "pad_token_id": 0,
+  "position_embedding_type": "absolute",
+  "torch_dtype": "float32",
+  "transformers_version": "4.52.4",
+  "type_vocab_size": 2,
+  "use_cache": true,
+  "vocab_size": 30522
+}

run-2/checkpoint-750/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:808582b9cc17178cc14e0f74f5d024452a9d04a7a4201ad7063501558f6cd336
+size 438050928

run-2/checkpoint-750/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:56a409175314ee85b05a73d63e6ff68bc3fb28b2439fa2007a50aa11821496a1
+size 876222970

run-2/checkpoint-750/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1b4f28eda59ed807ad79d8c5ee4c38e6dc5bdb88151d1f8b85343fdb7c170e46
+size 14244

run-2/checkpoint-750/scaler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:c81c4bed410b5f7275224ab7c747fbdb55beb1b2bd651a613603d5868aa5a2d9
+size 988

run-2/checkpoint-750/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5d975b57d7cb8d2ee1ed6615e320380089883d27823ebb20b6b01e4f67f42d8b
+size 1064

run-2/checkpoint-750/trainer_state.json ADDED Viewed

	@@ -0,0 +1,174 @@

+{
+  "best_global_step": 750,
+  "best_metric": 0.9368632136576173,
+  "best_model_checkpoint": "./results/run-2/checkpoint-750",
+  "epoch": 3.0,
+  "eval_steps": 500,
+  "global_step": 750,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 0.4,
+      "grad_norm": 1.9651439189910889,
+      "learning_rate": 1.386e-05,
+      "loss": 1.62,
+      "step": 100
+    },
+    {
+      "epoch": 0.8,
+      "grad_norm": 1.7124375104904175,
+      "learning_rate": 2.7859999999999998e-05,
+      "loss": 0.6715,
+      "step": 200
+    },
+    {
+      "epoch": 1.0,
+      "eval_clientelism_accuracy": 0.906,
+      "eval_clientelism_f1": 0.897770957127279,
+      "eval_discipline_among_poor_accuracy": 0.968,
+      "eval_discipline_among_poor_f1": 0.9615640823822642,
+      "eval_economic_policy_accuracy": 0.875,
+      "eval_economic_policy_f1": 0.860741831908876,
+      "eval_loss": 0.47158536314964294,
+      "eval_marcos_duterte_alliance_accuracy": 0.846,
+      "eval_marcos_duterte_alliance_f1": 0.839269429776376,
+      "eval_overall_accuracy": 0.8419375,
+      "eval_overall_f1": 0.8147702755950403,
+      "eval_populism_accuracy": 0.645,
+      "eval_populism_f1": 0.5639441230465639,
+      "eval_regionalism_accuracy": 0.963,
+      "eval_regionalism_f1": 0.9482946255564284,
+      "eval_runtime": 2.5786,
+      "eval_samples_per_second": 775.6,
+      "eval_security_accuracy": 0.8275,
+      "eval_security_f1": 0.8006718398296887,
+      "eval_steps_per_second": 48.475,
+      "eval_uniteam_positive_campaign_accuracy": 0.705,
+      "eval_uniteam_positive_campaign_f1": 0.645905315132846,
+      "step": 250
+    },
+    {
+      "epoch": 1.2,
+      "grad_norm": 1.057538390159607,
+      "learning_rate": 4.1859999999999996e-05,
+      "loss": 0.4761,
+      "step": 300
+    },
+    {
+      "epoch": 1.6,
+      "grad_norm": 1.3597612380981445,
+      "learning_rate": 5.586e-05,
+      "loss": 0.3424,
+      "step": 400
+    },
+    {
+      "epoch": 2.0,
+      "grad_norm": 1.4431722164154053,
+      "learning_rate": 6.986e-05,
+      "loss": 0.2992,
+      "step": 500
+    },
+    {
+      "epoch": 2.0,
+      "eval_clientelism_accuracy": 0.949,
+      "eval_clientelism_f1": 0.9452056048035566,
+      "eval_discipline_among_poor_accuracy": 0.975,
+      "eval_discipline_among_poor_f1": 0.9713251858526694,
+      "eval_economic_policy_accuracy": 0.946,
+      "eval_economic_policy_f1": 0.9447796444638326,
+      "eval_loss": 0.2847972512245178,
+      "eval_marcos_duterte_alliance_accuracy": 0.905,
+      "eval_marcos_duterte_alliance_f1": 0.8931059723613894,
+      "eval_overall_accuracy": 0.9166874999999999,
+      "eval_overall_f1": 0.9105488176533825,
+      "eval_populism_accuracy": 0.7845,
+      "eval_populism_f1": 0.7756216085131521,
+      "eval_regionalism_accuracy": 0.977,
+      "eval_regionalism_f1": 0.9741474164644311,
+      "eval_runtime": 2.6526,
+      "eval_samples_per_second": 753.967,
+      "eval_security_accuracy": 0.943,
+      "eval_security_f1": 0.934999506345517,
+      "eval_steps_per_second": 47.123,
+      "eval_uniteam_positive_campaign_accuracy": 0.854,
+      "eval_uniteam_positive_campaign_f1": 0.8452056024225119,
+      "step": 500
+    },
+    {
+      "epoch": 2.4,
+      "grad_norm": 1.2808622121810913,
+      "learning_rate": 6.0759999999999994e-05,
+      "loss": 0.2152,
+      "step": 600
+    },
+    {
+      "epoch": 2.8,
+      "grad_norm": 0.9503875970840454,
+      "learning_rate": 5.142666666666667e-05,
+      "loss": 0.1943,
+      "step": 700
+    },
+    {
+      "epoch": 3.0,
+      "eval_clientelism_accuracy": 0.954,
+      "eval_clientelism_f1": 0.9527297868951113,
+      "eval_discipline_among_poor_accuracy": 0.977,
+      "eval_discipline_among_poor_f1": 0.975000371636688,
+      "eval_economic_policy_accuracy": 0.9495,
+      "eval_economic_policy_f1": 0.9493495071116878,
+      "eval_loss": 0.2306518405675888,
+      "eval_marcos_duterte_alliance_accuracy": 0.9325,
+      "eval_marcos_duterte_alliance_f1": 0.9269038834173264,
+      "eval_overall_accuracy": 0.9380625,
+      "eval_overall_f1": 0.9368632136576173,
+      "eval_populism_accuracy": 0.8745,
+      "eval_populism_f1": 0.8748340797000174,
+      "eval_regionalism_accuracy": 0.9735,
+      "eval_regionalism_f1": 0.9749884932130427,
+      "eval_runtime": 2.7943,
+      "eval_samples_per_second": 715.749,
+      "eval_security_accuracy": 0.953,
+      "eval_security_f1": 0.9496585291837848,
+      "eval_steps_per_second": 44.734,
+      "eval_uniteam_positive_campaign_accuracy": 0.8905,
+      "eval_uniteam_positive_campaign_f1": 0.8914410581032797,
+      "step": 750
+    }
+  ],
+  "logging_steps": 100,
+  "max_steps": 1250,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 5,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "EarlyStoppingCallback": {
+      "args": {
+        "early_stopping_patience": 2,
+        "early_stopping_threshold": 0.0
+      },
+      "attributes": {
+        "early_stopping_patience_counter": 0
+      }
+    },
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": false
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 5895275151360000.0,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "gradient_accumulation_steps": 2,
+    "learning_rate": 7e-05,
+    "num_train_epochs": 5
+  }
+}

run-2/checkpoint-750/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:00db87b530f828066f2ef5374605ff7c067cf8bcaf6d07e1ad2162b4b005b6c2
+size 5368

run-3/checkpoint-125/config.json ADDED Viewed

	@@ -0,0 +1,93 @@

+{
+  "architectures": [
+    "BertForSequenceClassification"
+  ],
+  "attention_probs_dropout_prob": 0.1,
+  "classifier_dropout": null,
+  "gradient_checkpointing": false,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
+  "hidden_size": 768,
+  "id2label": {
+    "0": "LABEL_0",
+    "1": "LABEL_1",
+    "2": "LABEL_2",
+    "3": "LABEL_3",
+    "4": "LABEL_4",
+    "5": "LABEL_5",
+    "6": "LABEL_6",
+    "7": "LABEL_7",
+    "8": "LABEL_8",
+    "9": "LABEL_9",
+    "10": "LABEL_10",
+    "11": "LABEL_11",
+    "12": "LABEL_12",
+    "13": "LABEL_13",
+    "14": "LABEL_14",
+    "15": "LABEL_15",
+    "16": "LABEL_16",
+    "17": "LABEL_17",
+    "18": "LABEL_18",
+    "19": "LABEL_19",
+    "20": "LABEL_20",
+    "21": "LABEL_21",
+    "22": "LABEL_22",
+    "23": "LABEL_23",
+    "24": "LABEL_24",
+    "25": "LABEL_25",
+    "26": "LABEL_26",
+    "27": "LABEL_27",
+    "28": "LABEL_28",
+    "29": "LABEL_29",
+    "30": "LABEL_30",
+    "31": "LABEL_31"
+  },
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "label2id": {
+    "LABEL_0": 0,
+    "LABEL_1": 1,
+    "LABEL_10": 10,
+    "LABEL_11": 11,
+    "LABEL_12": 12,
+    "LABEL_13": 13,
+    "LABEL_14": 14,
+    "LABEL_15": 15,
+    "LABEL_16": 16,
+    "LABEL_17": 17,
+    "LABEL_18": 18,
+    "LABEL_19": 19,
+    "LABEL_2": 2,
+    "LABEL_20": 20,
+    "LABEL_21": 21,
+    "LABEL_22": 22,
+    "LABEL_23": 23,
+    "LABEL_24": 24,
+    "LABEL_25": 25,
+    "LABEL_26": 26,
+    "LABEL_27": 27,
+    "LABEL_28": 28,
+    "LABEL_29": 29,
+    "LABEL_3": 3,
+    "LABEL_30": 30,
+    "LABEL_31": 31,
+    "LABEL_4": 4,
+    "LABEL_5": 5,
+    "LABEL_6": 6,
+    "LABEL_7": 7,
+    "LABEL_8": 8,
+    "LABEL_9": 9
+  },
+  "layer_norm_eps": 1e-12,
+  "max_position_embeddings": 512,
+  "model_type": "bert",
+  "num_attention_heads": 12,
+  "num_hidden_layers": 12,
+  "pad_token_id": 0,
+  "position_embedding_type": "absolute",
+  "torch_dtype": "float32",
+  "transformers_version": "4.52.4",
+  "type_vocab_size": 2,
+  "use_cache": true,
+  "vocab_size": 30522
+}

run-3/checkpoint-125/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:a1eb23c4286bb70d60d37d169b4eb4cdef72b61f202122c30bc3dc70490c88d5
+size 438050928

run-3/checkpoint-125/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e9a736f2ba4a71589fdd75362155e09e41897aead5cdc665037e549e91a4bc42
+size 876222970

run-3/checkpoint-125/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:799803ad7ab7f8698f94ca574007f727508e835fc64b1230bf7233244f9b79ed
+size 14244

run-3/checkpoint-125/scaler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:09ec4bc8ac5689de95220cb2c918e80da7bc5dbd574ea051273889b34f7a4e18
+size 988

run-3/checkpoint-125/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:8cfe190de1a75396ad38ec62083e147ff07aaec92a9a25aa4684d42a688f26ae
+size 1064

run-3/checkpoint-125/trainer_state.json ADDED Viewed

	@@ -0,0 +1,80 @@

+{
+  "best_global_step": 125,
+  "best_metric": 0.7271188688911245,
+  "best_model_checkpoint": "./results/run-3/checkpoint-125",
+  "epoch": 1.0,
+  "eval_steps": 500,
+  "global_step": 125,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 0.8,
+      "grad_norm": 1.825743556022644,
+      "learning_rate": 9.900000000000002e-06,
+      "loss": 1.7201,
+      "step": 100
+    },
+    {
+      "epoch": 1.0,
+      "eval_clientelism_accuracy": 0.8205,
+      "eval_clientelism_f1": 0.7482918732782369,
+      "eval_discipline_among_poor_accuracy": 0.955,
+      "eval_discipline_among_poor_f1": 0.9330179028132993,
+      "eval_economic_policy_accuracy": 0.7925,
+      "eval_economic_policy_f1": 0.7007601115760111,
+      "eval_loss": 0.7271444797515869,
+      "eval_marcos_duterte_alliance_accuracy": 0.769,
+      "eval_marcos_duterte_alliance_f1": 0.7375199658801759,
+      "eval_overall_accuracy": 0.790125,
+      "eval_overall_f1": 0.7271188688911245,
+      "eval_populism_accuracy": 0.5955,
+      "eval_populism_f1": 0.5190889398077297,
+      "eval_regionalism_accuracy": 0.941,
+      "eval_regionalism_f1": 0.9123967027305513,
+      "eval_runtime": 2.4983,
+      "eval_samples_per_second": 800.535,
+      "eval_security_accuracy": 0.773,
+      "eval_security_f1": 0.6801353637901861,
+      "eval_steps_per_second": 50.033,
+      "eval_uniteam_positive_campaign_accuracy": 0.6745,
+      "eval_uniteam_positive_campaign_f1": 0.5857400912528054,
+      "step": 125
+    }
+  ],
+  "logging_steps": 100,
+  "max_steps": 625,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 5,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "EarlyStoppingCallback": {
+      "args": {
+        "early_stopping_patience": 2,
+        "early_stopping_threshold": 0.0
+      },
+      "attributes": {
+        "early_stopping_patience_counter": 0
+      }
+    },
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": false
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 1684364328960000.0,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "gradient_accumulation_steps": 4,
+    "learning_rate": 5e-05,
+    "num_train_epochs": 5
+  }
+}

run-3/checkpoint-125/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:548010b4f91c450431306186acf7c2d2640e58960634c399b136e2be17e8cbe9
+size 5368

training_args.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:00db87b530f828066f2ef5374605ff7c067cf8bcaf6d07e1ad2162b4b005b6c2
 size 5368

 version https://git-lfs.github.com/spec/v1
+oid sha256:548010b4f91c450431306186acf7c2d2640e58960634c399b136e2be17e8cbe9
 size 5368