wonderwind271 commited on Jun 13, 2025

Commit

12d0bc6

verified ·

1 Parent(s): be11506

vl-s42

Browse files

This view is limited to 50 files because it contains too many changes. See raw diff

Files changed (50) hide show

.gitattributes +25 -0
trabank_vl_dino_pretrain2/checkpoint-0/config.json +53 -0
trabank_vl_dino_pretrain2/checkpoint-0/generation_config.json +7 -0
trabank_vl_dino_pretrain2/checkpoint-0/model.safetensors +3 -0
trabank_vl_dino_pretrain2/checkpoint-0/special_tokens_map.json +24 -0
trabank_vl_dino_pretrain2/checkpoint-0/tokenizer_config.json +44 -0
trabank_vl_dino_pretrain2/checkpoint-0/training_args.bin +3 -0
trabank_vl_dino_pretrain2/checkpoint-0/vocab.json +0 -0
trabank_vl_dino_pretrain2/checkpoint-10000/config.json +53 -0
trabank_vl_dino_pretrain2/checkpoint-10000/generation_config.json +7 -0
trabank_vl_dino_pretrain2/checkpoint-10000/model.safetensors +3 -0
trabank_vl_dino_pretrain2/checkpoint-10000/optimizer.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-10000/rng_state_0.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-10000/rng_state_1.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-10000/scheduler.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-10000/special_tokens_map.json +24 -0
trabank_vl_dino_pretrain2/checkpoint-10000/tokenizer_config.json +44 -0
trabank_vl_dino_pretrain2/checkpoint-10000/trainer_state.json +0 -0
trabank_vl_dino_pretrain2/checkpoint-10000/training_args.bin +3 -0
trabank_vl_dino_pretrain2/checkpoint-10000/vocab.json +0 -0
trabank_vl_dino_pretrain2/checkpoint-100000/config.json +53 -0
trabank_vl_dino_pretrain2/checkpoint-100000/generation_config.json +7 -0
trabank_vl_dino_pretrain2/checkpoint-100000/model.safetensors +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/optimizer.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/rng_state_0.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/rng_state_1.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/scheduler.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/special_tokens_map.json +24 -0
trabank_vl_dino_pretrain2/checkpoint-100000/tokenizer_config.json +44 -0
trabank_vl_dino_pretrain2/checkpoint-100000/trainer_state.json +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/training_args.bin +3 -0
trabank_vl_dino_pretrain2/checkpoint-100000/vocab.json +0 -0
trabank_vl_dino_pretrain2/checkpoint-110000/config.json +53 -0
trabank_vl_dino_pretrain2/checkpoint-110000/generation_config.json +7 -0
trabank_vl_dino_pretrain2/checkpoint-110000/model.safetensors +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/optimizer.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/rng_state_0.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/rng_state_1.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/scheduler.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/special_tokens_map.json +24 -0
trabank_vl_dino_pretrain2/checkpoint-110000/tokenizer_config.json +44 -0
trabank_vl_dino_pretrain2/checkpoint-110000/trainer_state.json +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/training_args.bin +3 -0
trabank_vl_dino_pretrain2/checkpoint-110000/vocab.json +0 -0
trabank_vl_dino_pretrain2/checkpoint-120000/config.json +53 -0
trabank_vl_dino_pretrain2/checkpoint-120000/generation_config.json +7 -0
trabank_vl_dino_pretrain2/checkpoint-120000/model.safetensors +3 -0
trabank_vl_dino_pretrain2/checkpoint-120000/optimizer.pt +3 -0
trabank_vl_dino_pretrain2/checkpoint-120000/rng_state_0.pth +3 -0
trabank_vl_dino_pretrain2/checkpoint-120000/rng_state_1.pth +3 -0

.gitattributes CHANGED Viewed

@@ -33,3 +33,28 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
 *.zip filter=lfs diff=lfs merge=lfs -text
 *.zst filter=lfs diff=lfs merge=lfs -text
 *tfevents* filter=lfs diff=lfs merge=lfs -text

 *.zip filter=lfs diff=lfs merge=lfs -text
 *.zst filter=lfs diff=lfs merge=lfs -text
 *tfevents* filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-100000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-110000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-120000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-130000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-140000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-150000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-160000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-170000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-180000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-190000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-200000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-210000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-220000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-230000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-240000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-250000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-260000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-270000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-280000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-290000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-300000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-70000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-80000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/checkpoint-90000/trainer_state.json filter=lfs diff=lfs merge=lfs -text
+trabank_vl_dino_pretrain2/trainer_state.json filter=lfs diff=lfs merge=lfs -text

trabank_vl_dino_pretrain2/checkpoint-0/config.json ADDED Viewed

	@@ -0,0 +1,53 @@

+{
+  "activation_function": "gelu_new",
+  "architectures": [
+    "LlavaGPTForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "attn_pdrop": 0.1,
+  "bos_token_id": 1,
+  "detect_loss": false,
+  "embd_pdrop": 0.1,
+  "eos_token_id": 2,
+  "freeze_mm_mlp_adapter": false,
+  "hidden_act": "silu",
+  "image_aspect_ratio": "pad",
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": true,
+  "layer_norm_epsilon": 1e-05,
+  "mm_projector_lr": null,
+  "mm_use_im_patch_token": false,
+  "mm_use_im_start_end": false,
+  "mm_vision_tower": "dino",
+  "model_type": "gpt2",
+  "n_embd": 768,
+  "n_head": 12,
+  "n_inner": null,
+  "n_layer": 12,
+  "n_positions": 1024,
+  "num_key_value_heads": 12,
+  "reorder_and_upcast_attn": false,
+  "resid_pdrop": 0.1,
+  "rms_norm_eps": 1e-05,
+  "rope_scaling": null,
+  "rope_theta": 10000.0,
+  "scale_attn_by_inverse_layer_idx": false,
+  "scale_attn_weights": true,
+  "summary_activation": null,
+  "summary_first_dropout": 0.1,
+  "summary_proj_to_labels": true,
+  "summary_type": "cls_index",
+  "summary_use_proj": true,
+  "tie_word_embeddings": false,
+  "tokenizer_class": "LlamaTokenizer",
+  "tokenizer_model_max_length": 2048,
+  "tokenizer_padding_side": "right",
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.0",
+  "tune_mm_mlp_adapter": false,
+  "use_cache": false,
+  "vision_tower_type": "dino",
+  "vocab_size": 25005
+}

trabank_vl_dino_pretrain2/checkpoint-0/generation_config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "_from_model_config": true,
+  "bos_token_id": 1,
+  "eos_token_id": 2,
+  "transformers_version": "4.38.0",
+  "use_cache": false
+}

trabank_vl_dino_pretrain2/checkpoint-0/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:faa84e86541e3086ec685a513764937c1040a222727ce57d02a8e059fa87c929
+size 670210056

trabank_vl_dino_pretrain2/checkpoint-0/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,24 @@

+{
+  "bos_token": {
+    "content": "<SOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "eos_token": {
+    "content": "<EOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": "<UNK>",
+  "unk_token": {
+    "content": "<UNK>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

trabank_vl_dino_pretrain2/checkpoint-0/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,44 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "<PAD>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "1": {
+      "content": "<UNK>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "2": {
+      "content": "<SOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "3": {
+      "content": "<EOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "bos_token": "<SOS>",
+  "clean_up_tokenization_spaces": true,
+  "eos_token": "<EOS>",
+  "model_max_length": 2048,
+  "pad_token": "<UNK>",
+  "padding_side": "right",
+  "tokenizer_class": "CustomWordTokenizer",
+  "unk_token": "<UNK>"
+}

trabank_vl_dino_pretrain2/checkpoint-0/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:2e5cac6962892ba51458c4ff730d1530be81b09903879daeaa8214d165a7d5fb
+size 5905

trabank_vl_dino_pretrain2/checkpoint-0/vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff

trabank_vl_dino_pretrain2/checkpoint-10000/config.json ADDED Viewed

	@@ -0,0 +1,53 @@

+{
+  "activation_function": "gelu_new",
+  "architectures": [
+    "LlavaGPTForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "attn_pdrop": 0.1,
+  "bos_token_id": 1,
+  "detect_loss": false,
+  "embd_pdrop": 0.1,
+  "eos_token_id": 2,
+  "freeze_mm_mlp_adapter": false,
+  "hidden_act": "silu",
+  "image_aspect_ratio": "pad",
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": true,
+  "layer_norm_epsilon": 1e-05,
+  "mm_projector_lr": null,
+  "mm_use_im_patch_token": false,
+  "mm_use_im_start_end": false,
+  "mm_vision_tower": "dino",
+  "model_type": "gpt2",
+  "n_embd": 768,
+  "n_head": 12,
+  "n_inner": null,
+  "n_layer": 12,
+  "n_positions": 1024,
+  "num_key_value_heads": 12,
+  "reorder_and_upcast_attn": false,
+  "resid_pdrop": 0.1,
+  "rms_norm_eps": 1e-05,
+  "rope_scaling": null,
+  "rope_theta": 10000.0,
+  "scale_attn_by_inverse_layer_idx": false,
+  "scale_attn_weights": true,
+  "summary_activation": null,
+  "summary_first_dropout": 0.1,
+  "summary_proj_to_labels": true,
+  "summary_type": "cls_index",
+  "summary_use_proj": true,
+  "tie_word_embeddings": false,
+  "tokenizer_class": "LlamaTokenizer",
+  "tokenizer_model_max_length": 2048,
+  "tokenizer_padding_side": "right",
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.0",
+  "tune_mm_mlp_adapter": false,
+  "use_cache": false,
+  "vision_tower_type": "dino",
+  "vocab_size": 25005
+}

trabank_vl_dino_pretrain2/checkpoint-10000/generation_config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "_from_model_config": true,
+  "bos_token_id": 1,
+  "eos_token_id": 2,
+  "transformers_version": "4.38.0",
+  "use_cache": false
+}

trabank_vl_dino_pretrain2/checkpoint-10000/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:90ef9fa2011850327f84e64fb9db9c572d76cc4d952811b832f86fb7e3ff4a74
+size 670210056

trabank_vl_dino_pretrain2/checkpoint-10000/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1c5ae70950b05fea1c2935aff21b1e19e999b02fbe44575fcf3c2314d0afd8b9
+size 994126027

trabank_vl_dino_pretrain2/checkpoint-10000/rng_state_0.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e800e27c87e3293aa81fb875383cedc75fbc8052937e1473bed31266c559561f
+size 14917

trabank_vl_dino_pretrain2/checkpoint-10000/rng_state_1.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:638e8f22db8db88620f5f0208f57f1bcb0b2b1565c8b9c04bd676d85e2a3d1ef
+size 14917

trabank_vl_dino_pretrain2/checkpoint-10000/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4aeaedc1e6e4ac9ac4f28a8691779e4a906da7b051e36d0bbe385033f066178d
+size 1465

trabank_vl_dino_pretrain2/checkpoint-10000/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,24 @@

+{
+  "bos_token": {
+    "content": "<SOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "eos_token": {
+    "content": "<EOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": "<UNK>",
+  "unk_token": {
+    "content": "<UNK>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

trabank_vl_dino_pretrain2/checkpoint-10000/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,44 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "<PAD>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "1": {
+      "content": "<UNK>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "2": {
+      "content": "<SOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "3": {
+      "content": "<EOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "bos_token": "<SOS>",
+  "clean_up_tokenization_spaces": true,
+  "eos_token": "<EOS>",
+  "model_max_length": 2048,
+  "pad_token": "<UNK>",
+  "padding_side": "right",
+  "tokenizer_class": "CustomWordTokenizer",
+  "unk_token": "<UNK>"
+}

trabank_vl_dino_pretrain2/checkpoint-10000/trainer_state.json ADDED Viewed

The diff for this file is too large to render. See raw diff

trabank_vl_dino_pretrain2/checkpoint-10000/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:2d9afc445078d7a0e08b66aa2a6a92038125de135bbe646908f84d8a45adb76c
+size 5905

trabank_vl_dino_pretrain2/checkpoint-10000/vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff

trabank_vl_dino_pretrain2/checkpoint-100000/config.json ADDED Viewed

	@@ -0,0 +1,53 @@

+{
+  "activation_function": "gelu_new",
+  "architectures": [
+    "LlavaGPTForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "attn_pdrop": 0.1,
+  "bos_token_id": 1,
+  "detect_loss": false,
+  "embd_pdrop": 0.1,
+  "eos_token_id": 2,
+  "freeze_mm_mlp_adapter": false,
+  "hidden_act": "silu",
+  "image_aspect_ratio": "pad",
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": true,
+  "layer_norm_epsilon": 1e-05,
+  "mm_projector_lr": null,
+  "mm_use_im_patch_token": false,
+  "mm_use_im_start_end": false,
+  "mm_vision_tower": "dino",
+  "model_type": "gpt2",
+  "n_embd": 768,
+  "n_head": 12,
+  "n_inner": null,
+  "n_layer": 12,
+  "n_positions": 1024,
+  "num_key_value_heads": 12,
+  "reorder_and_upcast_attn": false,
+  "resid_pdrop": 0.1,
+  "rms_norm_eps": 1e-05,
+  "rope_scaling": null,
+  "rope_theta": 10000.0,
+  "scale_attn_by_inverse_layer_idx": false,
+  "scale_attn_weights": true,
+  "summary_activation": null,
+  "summary_first_dropout": 0.1,
+  "summary_proj_to_labels": true,
+  "summary_type": "cls_index",
+  "summary_use_proj": true,
+  "tie_word_embeddings": false,
+  "tokenizer_class": "LlamaTokenizer",
+  "tokenizer_model_max_length": 2048,
+  "tokenizer_padding_side": "right",
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.0",
+  "tune_mm_mlp_adapter": false,
+  "use_cache": false,
+  "vision_tower_type": "dino",
+  "vocab_size": 25005
+}

trabank_vl_dino_pretrain2/checkpoint-100000/generation_config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "_from_model_config": true,
+  "bos_token_id": 1,
+  "eos_token_id": 2,
+  "transformers_version": "4.38.0",
+  "use_cache": false
+}

trabank_vl_dino_pretrain2/checkpoint-100000/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:00ca1d52fafd1ebb7155f384954fca2382ad880b08882a469f26305989683c44
+size 670210056

trabank_vl_dino_pretrain2/checkpoint-100000/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:fc1375d754822e018db4d5360d8b8139166bb624a0810a42081c9def61875b3e
+size 994126027

trabank_vl_dino_pretrain2/checkpoint-100000/rng_state_0.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4ab5a99960a2fc56c2309ba16e43bd3d1513a48678b37cd8816fb12463dbde78
+size 14917

trabank_vl_dino_pretrain2/checkpoint-100000/rng_state_1.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ca4afb4b6b62eddff7750304b1a91e9b4ebc1efbbcca96cf8995080f0a278d3e
+size 14917

trabank_vl_dino_pretrain2/checkpoint-100000/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ae0a341cc06fce0d21bf7169c4964d449cefbf6ff943711fe1a78667c32ac1df
+size 1465

trabank_vl_dino_pretrain2/checkpoint-100000/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,24 @@

+{
+  "bos_token": {
+    "content": "<SOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "eos_token": {
+    "content": "<EOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": "<UNK>",
+  "unk_token": {
+    "content": "<UNK>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

trabank_vl_dino_pretrain2/checkpoint-100000/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,44 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "<PAD>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "1": {
+      "content": "<UNK>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "2": {
+      "content": "<SOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "3": {
+      "content": "<EOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "bos_token": "<SOS>",
+  "clean_up_tokenization_spaces": true,
+  "eos_token": "<EOS>",
+  "model_max_length": 2048,
+  "pad_token": "<UNK>",
+  "padding_side": "right",
+  "tokenizer_class": "CustomWordTokenizer",
+  "unk_token": "<UNK>"
+}

trabank_vl_dino_pretrain2/checkpoint-100000/trainer_state.json ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:80216fca02fe92705013c6dc436585ecffa30ba06861d7f7a95269b1c905c478
+size 16064021

trabank_vl_dino_pretrain2/checkpoint-100000/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:2d9afc445078d7a0e08b66aa2a6a92038125de135bbe646908f84d8a45adb76c
+size 5905

trabank_vl_dino_pretrain2/checkpoint-100000/vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff

trabank_vl_dino_pretrain2/checkpoint-110000/config.json ADDED Viewed

	@@ -0,0 +1,53 @@

+{
+  "activation_function": "gelu_new",
+  "architectures": [
+    "LlavaGPTForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "attn_pdrop": 0.1,
+  "bos_token_id": 1,
+  "detect_loss": false,
+  "embd_pdrop": 0.1,
+  "eos_token_id": 2,
+  "freeze_mm_mlp_adapter": false,
+  "hidden_act": "silu",
+  "image_aspect_ratio": "pad",
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": true,
+  "layer_norm_epsilon": 1e-05,
+  "mm_projector_lr": null,
+  "mm_use_im_patch_token": false,
+  "mm_use_im_start_end": false,
+  "mm_vision_tower": "dino",
+  "model_type": "gpt2",
+  "n_embd": 768,
+  "n_head": 12,
+  "n_inner": null,
+  "n_layer": 12,
+  "n_positions": 1024,
+  "num_key_value_heads": 12,
+  "reorder_and_upcast_attn": false,
+  "resid_pdrop": 0.1,
+  "rms_norm_eps": 1e-05,
+  "rope_scaling": null,
+  "rope_theta": 10000.0,
+  "scale_attn_by_inverse_layer_idx": false,
+  "scale_attn_weights": true,
+  "summary_activation": null,
+  "summary_first_dropout": 0.1,
+  "summary_proj_to_labels": true,
+  "summary_type": "cls_index",
+  "summary_use_proj": true,
+  "tie_word_embeddings": false,
+  "tokenizer_class": "LlamaTokenizer",
+  "tokenizer_model_max_length": 2048,
+  "tokenizer_padding_side": "right",
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.0",
+  "tune_mm_mlp_adapter": false,
+  "use_cache": false,
+  "vision_tower_type": "dino",
+  "vocab_size": 25005
+}

trabank_vl_dino_pretrain2/checkpoint-110000/generation_config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "_from_model_config": true,
+  "bos_token_id": 1,
+  "eos_token_id": 2,
+  "transformers_version": "4.38.0",
+  "use_cache": false
+}

trabank_vl_dino_pretrain2/checkpoint-110000/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:bd991a24cfcd991ee7a4cbdcf9c45f3ec1dd85b5e32b2b07349f4da70a02316f
+size 670210056

trabank_vl_dino_pretrain2/checkpoint-110000/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:0d727eb52ef6f54653b290538a3642a59995c84c9e5f0cd461c717d574b10d1d
+size 994126027

trabank_vl_dino_pretrain2/checkpoint-110000/rng_state_0.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:b9b970859d419f254ebd66499f32f0829a2c6ad668c2fa698c82a334b2ec740d
+size 14917

trabank_vl_dino_pretrain2/checkpoint-110000/rng_state_1.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d9591c12f58d7515ccb077714e9edbc56659e1783c5cec0bb5b4dde82022ded7
+size 14917

trabank_vl_dino_pretrain2/checkpoint-110000/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:db6067b09a272acaf5c85ca7f6a9e4741d56ce0b4071a298c7b399949ec03a20
+size 1465

trabank_vl_dino_pretrain2/checkpoint-110000/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,24 @@

+{
+  "bos_token": {
+    "content": "<SOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "eos_token": {
+    "content": "<EOS>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": "<UNK>",
+  "unk_token": {
+    "content": "<UNK>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

trabank_vl_dino_pretrain2/checkpoint-110000/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,44 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "<PAD>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "1": {
+      "content": "<UNK>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "2": {
+      "content": "<SOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "3": {
+      "content": "<EOS>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "bos_token": "<SOS>",
+  "clean_up_tokenization_spaces": true,
+  "eos_token": "<EOS>",
+  "model_max_length": 2048,
+  "pad_token": "<UNK>",
+  "padding_side": "right",
+  "tokenizer_class": "CustomWordTokenizer",
+  "unk_token": "<UNK>"
+}

trabank_vl_dino_pretrain2/checkpoint-110000/trainer_state.json ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:337d569e2828dd3fd086a52ffa81af64d92a6dc9f24b17cbfde25747b84c4bc9
+size 17683025

trabank_vl_dino_pretrain2/checkpoint-110000/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:2d9afc445078d7a0e08b66aa2a6a92038125de135bbe646908f84d8a45adb76c
+size 5905

trabank_vl_dino_pretrain2/checkpoint-110000/vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff

trabank_vl_dino_pretrain2/checkpoint-120000/config.json ADDED Viewed

	@@ -0,0 +1,53 @@

+{
+  "activation_function": "gelu_new",
+  "architectures": [
+    "LlavaGPTForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "attn_pdrop": 0.1,
+  "bos_token_id": 1,
+  "detect_loss": false,
+  "embd_pdrop": 0.1,
+  "eos_token_id": 2,
+  "freeze_mm_mlp_adapter": false,
+  "hidden_act": "silu",
+  "image_aspect_ratio": "pad",
+  "initializer_range": 0.02,
+  "intermediate_size": 3072,
+  "is_decoder": true,
+  "layer_norm_epsilon": 1e-05,
+  "mm_projector_lr": null,
+  "mm_use_im_patch_token": false,
+  "mm_use_im_start_end": false,
+  "mm_vision_tower": "dino",
+  "model_type": "gpt2",
+  "n_embd": 768,
+  "n_head": 12,
+  "n_inner": null,
+  "n_layer": 12,
+  "n_positions": 1024,
+  "num_key_value_heads": 12,
+  "reorder_and_upcast_attn": false,
+  "resid_pdrop": 0.1,
+  "rms_norm_eps": 1e-05,
+  "rope_scaling": null,
+  "rope_theta": 10000.0,
+  "scale_attn_by_inverse_layer_idx": false,
+  "scale_attn_weights": true,
+  "summary_activation": null,
+  "summary_first_dropout": 0.1,
+  "summary_proj_to_labels": true,
+  "summary_type": "cls_index",
+  "summary_use_proj": true,
+  "tie_word_embeddings": false,
+  "tokenizer_class": "LlamaTokenizer",
+  "tokenizer_model_max_length": 2048,
+  "tokenizer_padding_side": "right",
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.0",
+  "tune_mm_mlp_adapter": false,
+  "use_cache": false,
+  "vision_tower_type": "dino",
+  "vocab_size": 25005
+}

trabank_vl_dino_pretrain2/checkpoint-120000/generation_config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "_from_model_config": true,
+  "bos_token_id": 1,
+  "eos_token_id": 2,
+  "transformers_version": "4.38.0",
+  "use_cache": false
+}

trabank_vl_dino_pretrain2/checkpoint-120000/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d45552ece60140f6e63778efc55f311bbd80fc80c89272cd86016f90db301d67
+size 670210056

trabank_vl_dino_pretrain2/checkpoint-120000/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f645138e83688a85932b1e92166f38000c3cf17d7ad9d6a1b00ac2ddecec8288
+size 994126027

trabank_vl_dino_pretrain2/checkpoint-120000/rng_state_0.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:50b0d8fa412a6cc7161325ac315c1e403a306eb94fd435484bdf388232dd33d7
+size 14917

trabank_vl_dino_pretrain2/checkpoint-120000/rng_state_1.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d7304e2a27c7dba656585f92022f096860b4dc88b1503d5dcb71e348722b3a5c
+size 14917