Upload folder using huggingface_hub

Browse files

Files changed (12) hide show

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/config.json +22 -66
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/model.safetensors +2 -2
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/optimizer.pt +2 -2
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/rng_state.pth +1 -1
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/tokenizer.json +0 -0
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/tokenizer_config.json +5 -7
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/trainer_state.json +11 -2
open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/training_args.bin +1 -1
open-pii-masking-500k-ai4privacy/TokenBased-CRF/config.json +22 -66
open-pii-masking-500k-ai4privacy/TokenBased-CRF/pytorch_model.bin +2 -2
open-pii-masking-500k-ai4privacy/TokenBased-CRF/tokenizer.json +0 -0
open-pii-masking-500k-ai4privacy/TokenBased-CRF/tokenizer_config.json +5 -7

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/config.json CHANGED Viewed

@@ -1,23 +1,16 @@
 {
   "architectures": [
     "TransformerCrfForTokenClassification"
   ],
-  "attention_bias": false,
-  "attention_dropout": 0.0,
-  "bos_token_id": 50281,
-  "classifier_activation": "gelu",
-  "classifier_bias": false,
-  "classifier_dropout": 0.0,
-  "classifier_pooling": "mean",
-  "cls_token_id": 50281,
-  "decoder_bias": true,
-  "deterministic_flash_attn": false,
   "dtype": "float32",
-  "embedding_dropout": 0.0,
-  "eos_token_id": 50282,
-  "global_attn_every_n_layers": 3,
-  "gradient_checkpointing": false,
-  "hidden_activation": "gelu",
   "hidden_size": 768,
   "id2label": {
     "0": "O",
@@ -62,9 +55,9 @@
     "39": "B-ZIPCODE",
     "40": "I-ZIPCODE"
   },
-  "initializer_cutoff_factor": 2.0,
   "initializer_range": 0.02,
-  "intermediate_size": 1152,
   "label2id": {
     "B-AGE": 1,
     "B-BUILDINGNUM": 3,
@@ -108,57 +101,20 @@
     "I-ZIPCODE": 40,
     "O": 0
   },
-  "layer_norm_eps": 1e-05,
-  "layer_types": [
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention"
-  ],
-  "local_attention": 128,
-  "max_position_embeddings": 8192,
-  "mlp_bias": false,
-  "mlp_dropout": 0.0,
-  "model_type": "modernbert",
-  "norm_bias": false,
-  "norm_eps": 1e-05,
   "num_attention_heads": 12,
-  "num_hidden_layers": 22,
-  "pad_token_id": 50283,
-  "position_embedding_type": "absolute",
-  "rope_parameters": {
-    "full_attention": {
-      "rope_theta": 160000.0,
-      "rope_type": "default"
-    },
-    "sliding_attention": {
-      "rope_theta": 10000.0,
-      "rope_type": "default"
-    }
-  },
-  "sep_token_id": 50282,
-  "sparse_pred_ignore_index": -100,
-  "sparse_prediction": false,
   "tie_word_embeddings": true,
   "transformers_version": "5.3.0",
   "use_cache": false,
-  "vocab_size": 50368
 }

 {
+  "add_cross_attention": false,
   "architectures": [
     "TransformerCrfForTokenClassification"
   ],
+  "attention_probs_dropout_prob": 0.1,
+  "bos_token_id": null,
+  "classifier_dropout": null,
+  "directionality": "bidi",
   "dtype": "float32",
+  "eos_token_id": null,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
   "hidden_size": 768,
   "id2label": {
     "0": "O",
     "39": "B-ZIPCODE",
     "40": "I-ZIPCODE"
   },
   "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": false,
   "label2id": {
     "B-AGE": 1,
     "B-BUILDINGNUM": 3,
     "I-ZIPCODE": 40,
     "O": 0
   },
+  "layer_norm_eps": 1e-12,
+  "max_position_embeddings": 512,
+  "model_type": "bert",
   "num_attention_heads": 12,
+  "num_hidden_layers": 12,
+  "pad_token_id": 0,
+  "pooler_fc_size": 768,
+  "pooler_num_attention_heads": 12,
+  "pooler_num_fc_layers": 3,
+  "pooler_size_per_head": 128,
+  "pooler_type": "first_token_transform",
   "tie_word_embeddings": true,
   "transformers_version": "5.3.0",
+  "type_vocab_size": 2,
   "use_cache": false,
+  "vocab_size": 119547
 }

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/model.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:d14174801d86131d4e9aa1e0d6df08bdee9da59d4f0229589f19f0fc8c665319
-size 596205320

 version https://git-lfs.github.com/spec/v1
+oid sha256:988fe44477a3633c01f74398d42a45d67e46e272c89bf183fa404ff371ae3be6
+size 711572104

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/optimizer.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:bf89658473422e1c0214f65abe3aa73b5dbb589a8625dfca0af4d182e4eb0459
-size 1192499723

 version https://git-lfs.github.com/spec/v1
+oid sha256:6ce6f82e3963c0b944789c4de89f4bc3affccf18537a788f4d88f11ad232af03
+size 1418541131

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/rng_state.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:7991aa545aa2a9819a120cc19b48bfbf88dcaec121a00f803b9fbabf38d57d55
 size 14645

 version https://git-lfs.github.com/spec/v1
+oid sha256:f39eab8ed980549bfffcd8b948e8852bef979d820a1c47a898b2c9f270cc3986
 size 14645

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/tokenizer.json CHANGED Viewed

The diff for this file is too large to render. See raw diff

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/tokenizer_config.json CHANGED Viewed

@@ -1,17 +1,15 @@
 {
   "add_prefix_space": true,
   "backend": "tokenizers",
-  "clean_up_tokenization_spaces": true,
   "cls_token": "[CLS]",
   "is_local": false,
   "mask_token": "[MASK]",
-  "model_input_names": [
-    "input_ids",
-    "attention_mask"
-  ],
-  "model_max_length": 8192,
   "pad_token": "[PAD]",
   "sep_token": "[SEP]",
-  "tokenizer_class": "TokenizersBackend",
   "unk_token": "[UNK]"
 }

 {
   "add_prefix_space": true,
   "backend": "tokenizers",
   "cls_token": "[CLS]",
+  "do_lower_case": false,
   "is_local": false,
   "mask_token": "[MASK]",
+  "model_max_length": 512,
   "pad_token": "[PAD]",
   "sep_token": "[SEP]",
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "BertTokenizer",
   "unk_token": "[UNK]"
 }

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/trainer_state.json CHANGED Viewed

@@ -8,7 +8,16 @@
   "is_hyper_param_search": false,
   "is_local_process_zero": true,
   "is_world_process_zero": true,
-  "log_history": [],
   "logging_steps": 500,
   "max_steps": 63,
   "num_input_tokens_seen": 0,
@@ -26,7 +35,7 @@
       "attributes": {}
     }
   },
-  "total_flos": 169520547840000.0,
   "train_batch_size": 8,
   "trial_name": null,
   "trial_params": null

   "is_hyper_param_search": false,
   "is_local_process_zero": true,
   "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_loss": 24.712289810180664,
+      "eval_runtime": 1.6052,
+      "eval_samples_per_second": 62.296,
+      "eval_steps_per_second": 8.098,
+      "step": 63
+    }
+  ],
   "logging_steps": 500,
   "max_steps": 63,
   "num_input_tokens_seen": 0,
       "attributes": {}
     }
   },
+  "total_flos": 131604301824000.0,
   "train_batch_size": 8,
   "trial_name": null,
   "trial_params": null

open-pii-masking-500k-ai4privacy/TokenBased-CRF/checkpoint-63/training_args.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:4e5cf459979c7c595ed5f7c36736caa950a566433ec0f1908e8fc19c60ff5a21
 size 5201

 version https://git-lfs.github.com/spec/v1
+oid sha256:2b097dce68df0c542b37b2de4a187ae721b6306bf055ab3a17291f74d1ca8e64
 size 5201

open-pii-masking-500k-ai4privacy/TokenBased-CRF/config.json CHANGED Viewed

@@ -1,23 +1,16 @@
 {
   "architectures": [
     "TransformerCrfForTokenClassification"
   ],
-  "attention_bias": false,
-  "attention_dropout": 0.0,
-  "bos_token_id": 50281,
-  "classifier_activation": "gelu",
-  "classifier_bias": false,
-  "classifier_dropout": 0.0,
-  "classifier_pooling": "mean",
-  "cls_token_id": 50281,
-  "decoder_bias": true,
-  "deterministic_flash_attn": false,
   "dtype": "float32",
-  "embedding_dropout": 0.0,
-  "eos_token_id": 50282,
-  "global_attn_every_n_layers": 3,
-  "gradient_checkpointing": false,
-  "hidden_activation": "gelu",
   "hidden_size": 768,
   "id2label": {
     "0": "O",
@@ -62,9 +55,9 @@
     "39": "B-ZIPCODE",
     "40": "I-ZIPCODE"
   },
-  "initializer_cutoff_factor": 2.0,
   "initializer_range": 0.02,
-  "intermediate_size": 1152,
   "label2id": {
     "B-AGE": 1,
     "B-BUILDINGNUM": 3,
@@ -108,57 +101,20 @@
     "I-ZIPCODE": 40,
     "O": 0
   },
-  "layer_norm_eps": 1e-05,
-  "layer_types": [
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention",
-    "sliding_attention",
-    "sliding_attention",
-    "full_attention"
-  ],
-  "local_attention": 128,
-  "max_position_embeddings": 8192,
-  "mlp_bias": false,
-  "mlp_dropout": 0.0,
-  "model_type": "modernbert",
-  "norm_bias": false,
-  "norm_eps": 1e-05,
   "num_attention_heads": 12,
-  "num_hidden_layers": 22,
-  "pad_token_id": 50283,
-  "position_embedding_type": "absolute",
-  "rope_parameters": {
-    "full_attention": {
-      "rope_theta": 160000.0,
-      "rope_type": "default"
-    },
-    "sliding_attention": {
-      "rope_theta": 10000.0,
-      "rope_type": "default"
-    }
-  },
-  "sep_token_id": 50282,
-  "sparse_pred_ignore_index": -100,
-  "sparse_prediction": false,
   "tie_word_embeddings": true,
   "transformers_version": "5.3.0",
   "use_cache": false,
-  "vocab_size": 50368
 }

 {
+  "add_cross_attention": false,
   "architectures": [
     "TransformerCrfForTokenClassification"
   ],
+  "attention_probs_dropout_prob": 0.1,
+  "bos_token_id": null,
+  "classifier_dropout": null,
+  "directionality": "bidi",
   "dtype": "float32",
+  "eos_token_id": null,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
   "hidden_size": 768,
   "id2label": {
     "0": "O",
     "39": "B-ZIPCODE",
     "40": "I-ZIPCODE"
   },
   "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": false,
   "label2id": {
     "B-AGE": 1,
     "B-BUILDINGNUM": 3,
     "I-ZIPCODE": 40,
     "O": 0
   },
+  "layer_norm_eps": 1e-12,
+  "max_position_embeddings": 512,
+  "model_type": "bert",
   "num_attention_heads": 12,
+  "num_hidden_layers": 12,
+  "pad_token_id": 0,
+  "pooler_fc_size": 768,
+  "pooler_num_attention_heads": 12,
+  "pooler_num_fc_layers": 3,
+  "pooler_size_per_head": 128,
+  "pooler_type": "first_token_transform",
   "tie_word_embeddings": true,
   "transformers_version": "5.3.0",
+  "type_vocab_size": 2,
   "use_cache": false,
+  "vocab_size": 119547
 }

open-pii-masking-500k-ai4privacy/TokenBased-CRF/pytorch_model.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:66c257c667c083dd911145530aa329c1ffd0eea0f9dee2d51a4f92ba73d3ed70
-size 596251343

 version https://git-lfs.github.com/spec/v1
+oid sha256:4a2291ae713e208e7f762adbcbf6f6927ec7eb8be2dbeef176a6b54423d8c255
+size 711634067

open-pii-masking-500k-ai4privacy/TokenBased-CRF/tokenizer.json CHANGED Viewed

The diff for this file is too large to render. See raw diff

open-pii-masking-500k-ai4privacy/TokenBased-CRF/tokenizer_config.json CHANGED Viewed

@@ -1,17 +1,15 @@
 {
   "add_prefix_space": true,
   "backend": "tokenizers",
-  "clean_up_tokenization_spaces": true,
   "cls_token": "[CLS]",
   "is_local": false,
   "mask_token": "[MASK]",
-  "model_input_names": [
-    "input_ids",
-    "attention_mask"
-  ],
-  "model_max_length": 8192,
   "pad_token": "[PAD]",
   "sep_token": "[SEP]",
-  "tokenizer_class": "TokenizersBackend",
   "unk_token": "[UNK]"
 }

 {
   "add_prefix_space": true,
   "backend": "tokenizers",
   "cls_token": "[CLS]",
+  "do_lower_case": false,
   "is_local": false,
   "mask_token": "[MASK]",
+  "model_max_length": 512,
   "pad_token": "[PAD]",
   "sep_token": "[SEP]",
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "BertTokenizer",
   "unk_token": "[UNK]"
 }