bruhzair commited on
Commit
54c1a53
·
verified ·
1 Parent(s): ff565c1

Upload folder using huggingface_hub

Browse files
This view is limited to 50 files because it contains too many changes.   See raw diff
Files changed (50) hide show
  1. .gitattributes +1 -0
  2. README.md +47 -0
  3. chat_template.jinja +7 -0
  4. config.json +35 -0
  5. mergekit_config.yml +18 -0
  6. model-00001-of-00063.safetensors +3 -0
  7. model-00002-of-00063.safetensors +3 -0
  8. model-00003-of-00063.safetensors +3 -0
  9. model-00004-of-00063.safetensors +3 -0
  10. model-00005-of-00063.safetensors +3 -0
  11. model-00006-of-00063.safetensors +3 -0
  12. model-00007-of-00063.safetensors +3 -0
  13. model-00008-of-00063.safetensors +3 -0
  14. model-00009-of-00063.safetensors +3 -0
  15. model-00010-of-00063.safetensors +3 -0
  16. model-00011-of-00063.safetensors +3 -0
  17. model-00012-of-00063.safetensors +3 -0
  18. model-00013-of-00063.safetensors +3 -0
  19. model-00014-of-00063.safetensors +3 -0
  20. model-00015-of-00063.safetensors +3 -0
  21. model-00016-of-00063.safetensors +3 -0
  22. model-00017-of-00063.safetensors +3 -0
  23. model-00018-of-00063.safetensors +3 -0
  24. model-00019-of-00063.safetensors +3 -0
  25. model-00020-of-00063.safetensors +3 -0
  26. model-00021-of-00063.safetensors +3 -0
  27. model-00022-of-00063.safetensors +3 -0
  28. model-00023-of-00063.safetensors +3 -0
  29. model-00024-of-00063.safetensors +3 -0
  30. model-00025-of-00063.safetensors +3 -0
  31. model-00026-of-00063.safetensors +3 -0
  32. model-00027-of-00063.safetensors +3 -0
  33. model-00028-of-00063.safetensors +3 -0
  34. model-00029-of-00063.safetensors +3 -0
  35. model-00030-of-00063.safetensors +3 -0
  36. model-00031-of-00063.safetensors +3 -0
  37. model-00032-of-00063.safetensors +3 -0
  38. model-00033-of-00063.safetensors +3 -0
  39. model-00034-of-00063.safetensors +3 -0
  40. model-00035-of-00063.safetensors +3 -0
  41. model-00036-of-00063.safetensors +3 -0
  42. model-00037-of-00063.safetensors +3 -0
  43. model-00038-of-00063.safetensors +3 -0
  44. model-00039-of-00063.safetensors +3 -0
  45. model-00040-of-00063.safetensors +3 -0
  46. model-00041-of-00063.safetensors +3 -0
  47. model-00042-of-00063.safetensors +3 -0
  48. model-00043-of-00063.safetensors +3 -0
  49. model-00044-of-00063.safetensors +3 -0
  50. model-00045-of-00063.safetensors +3 -0
.gitattributes CHANGED
@@ -33,3 +33,4 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
 
 
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
36
+ tokenizer.json filter=lfs diff=lfs merge=lfs -text
README.md ADDED
@@ -0,0 +1,47 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ ---
2
+ base_model: []
3
+ library_name: transformers
4
+ tags:
5
+ - mergekit
6
+ - merge
7
+
8
+ ---
9
+ # prototype-0.4x329
10
+
11
+ This is a merge of pre-trained language models created using [mergekit](https://github.com/cg123/mergekit).
12
+
13
+ ## Merge Details
14
+ ### Merge Method
15
+
16
+ This model was merged using the [Multi-SLERP](https://goddard.blog/posts/multislerp-wow-what-a-cool-idea) merge method using /workspace/cache/models--deepcogito--cogito-v2-preview-llama-70B/snapshots/1e1d12e8eaebd6084a8dcf45ecdeaa2f4b8879ce as a base.
17
+
18
+ ### Models Merged
19
+
20
+ The following models were included in the merge:
21
+ * /workspace/cache/models--TheDrummer--Anubis-70B-v1/snapshots/e50d699bf6c21afcf4dbd9a8b4f73511b0366efb
22
+ * /workspace/cache/models--TheDrummer--Anubis-70B-v1.1/snapshots/47ea1a3368e8d161b09acbc8c211ba4212e4b466
23
+
24
+ ### Configuration
25
+
26
+ The following YAML configuration was used to produce this model:
27
+
28
+ ```yaml
29
+ models:
30
+ - model: /workspace/cache/models--TheDrummer--Anubis-70B-v1/snapshots/e50d699bf6c21afcf4dbd9a8b4f73511b0366efb
31
+ parameters:
32
+ weight: [0.5]
33
+ - model: /workspace/cache/models--TheDrummer--Anubis-70B-v1.1/snapshots/47ea1a3368e8d161b09acbc8c211ba4212e4b466
34
+ parameters:
35
+ weight: [0.5]
36
+ base_model: /workspace/cache/models--deepcogito--cogito-v2-preview-llama-70B/snapshots/1e1d12e8eaebd6084a8dcf45ecdeaa2f4b8879ce
37
+ merge_method: multislerp
38
+ tokenizer:
39
+ source: base
40
+ chat_template: llama3
41
+ parameters:
42
+ normalize_weights: false
43
+ eps: 1e-9
44
+ pad_to_multiple_of: 8
45
+ int8_mask: true
46
+ dtype: float32
47
+ ```
chat_template.jinja ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ {% set loop_messages = messages %}
2
+ {% for message in loop_messages %}
3
+ {% set content = '<|start_header_id|>' + message['role'] + '<|end_header_id|>\n\n'+ message['content'] | trim + '<|eot_id|>' %}
4
+ {% if loop.index0 == 0 %}{% set content = bos_token + content %}{% endif %}
5
+ {{ content }}
6
+ {% endfor %}
7
+ {% if add_generation_prompt %}{{ '<|start_header_id|>assistant<|end_header_id|>\n\n' }}{% endif %}
config.json ADDED
@@ -0,0 +1,35 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "architectures": [
3
+ "LlamaForCausalLM"
4
+ ],
5
+ "attention_bias": false,
6
+ "attention_dropout": 0.0,
7
+ "bos_token_id": 128000,
8
+ "eos_token_id": 128001,
9
+ "head_dim": 128,
10
+ "hidden_act": "silu",
11
+ "hidden_size": 8192,
12
+ "initializer_range": 0.02,
13
+ "intermediate_size": 28672,
14
+ "max_position_embeddings": 131072,
15
+ "mlp_bias": false,
16
+ "model_type": "llama",
17
+ "num_attention_heads": 64,
18
+ "num_hidden_layers": 80,
19
+ "num_key_value_heads": 8,
20
+ "pretraining_tp": 1,
21
+ "rms_norm_eps": 1e-05,
22
+ "rope_scaling": {
23
+ "factor": 8.0,
24
+ "high_freq_factor": 4.0,
25
+ "low_freq_factor": 1.0,
26
+ "original_max_position_embeddings": 8192,
27
+ "rope_type": "llama3"
28
+ },
29
+ "rope_theta": 500000.0,
30
+ "tie_word_embeddings": false,
31
+ "torch_dtype": "float32",
32
+ "transformers_version": "4.55.2",
33
+ "use_cache": true,
34
+ "vocab_size": 128256
35
+ }
mergekit_config.yml ADDED
@@ -0,0 +1,18 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ models:
2
+ - model: /workspace/cache/models--TheDrummer--Anubis-70B-v1/snapshots/e50d699bf6c21afcf4dbd9a8b4f73511b0366efb
3
+ parameters:
4
+ weight: [0.5]
5
+ - model: /workspace/cache/models--TheDrummer--Anubis-70B-v1.1/snapshots/47ea1a3368e8d161b09acbc8c211ba4212e4b466
6
+ parameters:
7
+ weight: [0.5]
8
+ base_model: /workspace/cache/models--deepcogito--cogito-v2-preview-llama-70B/snapshots/1e1d12e8eaebd6084a8dcf45ecdeaa2f4b8879ce
9
+ merge_method: multislerp
10
+ tokenizer:
11
+ source: base
12
+ chat_template: llama3
13
+ parameters:
14
+ normalize_weights: false
15
+ eps: 1e-9
16
+ pad_to_multiple_of: 8
17
+ int8_mask: true
18
+ dtype: float32
model-00001-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:1df450f7d7af904e8256bfad0042b066a6942ea4ee9533d0847e41032f825896
3
+ size 4362175752
model-00002-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:155e8dd4bda50386d554dd99cd8bee81493aefb7e49cdbc6744d33eedf39efb2
3
+ size 4429251992
model-00003-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:dd1ec20f4a7eaafce6a107b37452ffd37eeed79d539436604072511453a551f9
3
+ size 3154150224
model-00004-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:03ae15dae370db122e091e490a12256c2846e6ba5d6ddb36d068b5aa18a9c284
3
+ size 4236247288
model-00005-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:587fe359a27e0f2255010b107d696b7e920c09fa02093ae0020e877d81197eda
3
+ size 4462806536
model-00006-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:22e85521bf7f7d72396db1ad480d9181930bf7aa5e4b13e82361fd85f7bca9ad
3
+ size 2181071464
model-00007-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a81afecb58ae1935c206f6c480165f4c682caee2dc22e6589a72c40d1cf3ae0f
3
+ size 4504682872
model-00008-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5c1b49a1fb31f575a7bc9c78dc844d75298f0067e3b6a1509acd213b29a9b21d
3
+ size 4362208640
model-00009-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:bbea6b6ce8ccf3e0afa5a82b1cfd76ad3a76ec50e47db654273cda0f6ea85d6e
3
+ size 4630578456
model-00010-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5a27a4f35360f0635d954653a1281f17766ce6cc94fd3431b703b480e736d05b
3
+ size 4697687552
model-00011-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7538d227c6e0b6c2dbe859541e9a5c890db2cd8f0ed2ca5e42ebcf4b7fad497e
3
+ size 4362208640
model-00012-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2ab1cf8415bad5e25d857d80780cfe9e06cd92785870fc8d13a7157963f6059c
3
+ size 4362142888
model-00013-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:20418420db6217012bd860052de05d11d93cfb757f23eaebb83f59efdb69bd5a
3
+ size 4966123128
model-00014-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:efbd233693793fda525781ad22399cfd1819bbbe7f5f2010aa2aad93ed498bbb
3
+ size 4362208640
model-00015-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:693d73d72db36f2db3206c80e1c84c05a3ccdae8ba28d96eccc95eec06e29d19
3
+ size 4362142888
model-00016-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:32f376bb9676e7f340b95a5a7c2620068111fb6f52328be7c45e5d589b2ef42a
3
+ size 4966123128
model-00017-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0f50394df80aab5d2646d1a2e2b487bacbcfa7a1b61816adcb7cacba2a901df0
3
+ size 4362208640
model-00018-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0d7e767c14e514b3666bfc36a2bf26fd2558a968b8e0d35ac4138790da611f19
3
+ size 4362142888
model-00019-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7ea1da2fadf94d6e3ed5e0e03f6b33c651ebd8db356bfcd4fa90961d459e2346
3
+ size 4966156000
model-00020-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3f56fcd8c2a82286682ba2e1230e6e49de3ba90d0766702b82603709edd4c409
3
+ size 4362142888
model-00021-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d077cd518bec2be8ad87742662890c40f27b5f90fcb081ff5700293adb2fdb0e
3
+ size 4395730320
model-00022-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0a37e7a9e1bb902350b8a7f55425c3e0e80168b10f55e95538159da335f1c07a
3
+ size 4932601456
model-00023-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:29929b81894eec197b08fb983260d2acba87db92f1b2d435f45f9bb9e2f3e290
3
+ size 4395697432
model-00024-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c7483e79523cfad47a1bf32b0ed6a1076e414ac1a3f926332ee444d036279bf1
3
+ size 4395730320
model-00025-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:dd938a352a74ed5ef989ceb3e989a66e534f1d95681539809cc9a406f4963469
3
+ size 4630578456
model-00026-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:71d4d20dc1c70693c7a2949579185a0354b72db3ace129b25ab6fd81ba4a4014
3
+ size 4664165888
model-00027-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ac15da2ab3f920f6d1c0d627d31c32ab08543a4fefa7932d345de1488528a53b
3
+ size 4395730320
model-00028-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a39a84ff4a8df4049aa8644b5e66740b0f47012d334603f3195ac81d8b430dbe
3
+ size 4630578456
model-00029-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2b4925d11c3d1f7cd1dc823b8b66490a7db2c03a912d5a6392e58f0177f09aed
3
+ size 4664165888
model-00030-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:083babcd71b100daffcf84eed694b168d6d6a673028f12ccdfff0c999125c7be
3
+ size 4395730320
model-00031-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3fae1fcb82ee6803aaacff99f2b6013ec8de2cb6d95fc2e50c521a7d364e1528
3
+ size 4630578456
model-00032-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e12577594781a72274c3a56f7a426ff5a09286c64d97e0af1ba2d248f5d10b17
3
+ size 4664165880
model-00033-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:1b4d95672d0f7f841cdf2adb7839ae2727dd9556380a98589456ff9fe00db812
3
+ size 4395730320
model-00034-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:32e30b74b86263eeddaaf00fe724b93f2393229994c308add2bff5cdfa01c3b9
3
+ size 4630578456
model-00035-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0bdb28a9420a8439a637939cf0454442a2a8053bfb559025a03e46ca624d10dc
3
+ size 4664165880
model-00036-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:25f15b387e888a779cca0e0abdfb5a24f2e2a5a7d5eaa4f542740dc6d684eb2c
3
+ size 4395730320
model-00037-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:88144e23c401053e53cb95fa09e37f23b12e13e8f9e42f314088d1699857a3fb
3
+ size 4899014032
model-00038-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:65612dcf558fd575dbefa93891f9bc782adb84b5f99470196d6b2faed19ac34d
3
+ size 4395730304
model-00039-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a9a0a963a21b94daf5a38a724d94760d5d0c9b158ce043b3c9e67511e30a82a4
3
+ size 4395730320
model-00040-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:65ab29dd072503a9b582bf9a610d3b1986a3b087f248129ae02b90d77e282f65
3
+ size 4899014032
model-00041-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:328992287a9a3994ce1243a2049c17106556fe4b44a55c8450699f059f43e855
3
+ size 4395730304
model-00042-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:389abb00562ae9598e66a239f49548cf9f889e56326359b07307004fa9f8ed71
3
+ size 4395730320
model-00043-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:13f7e7aac7235b47d6b8574e777e302dddcdce13b66180a60c552dce71c7a38e
3
+ size 4630578456
model-00044-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:393036a3bb94c6f552ee55727f06ae6e5941088ab0a6f979011ec52c77c8166c
3
+ size 4664165880
model-00045-of-00063.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:73c2c4bae0ed4b9f90c8b828921429d4de3b2fee1372098e1dbb39c2b0f48e64
3
+ size 4395730320