Upload folder using huggingface_hub
Browse filesThis view is limited to 50 files because it contains too many changes.
See raw diff
- README.md +41 -0
- config.json +27 -0
- mergekit_config.yml +10 -0
- model-00001-of-00071.safetensors +3 -0
- model-00002-of-00071.safetensors +3 -0
- model-00003-of-00071.safetensors +3 -0
- model-00004-of-00071.safetensors +3 -0
- model-00005-of-00071.safetensors +3 -0
- model-00006-of-00071.safetensors +3 -0
- model-00007-of-00071.safetensors +3 -0
- model-00008-of-00071.safetensors +3 -0
- model-00009-of-00071.safetensors +3 -0
- model-00010-of-00071.safetensors +3 -0
- model-00011-of-00071.safetensors +3 -0
- model-00012-of-00071.safetensors +3 -0
- model-00013-of-00071.safetensors +3 -0
- model-00014-of-00071.safetensors +3 -0
- model-00015-of-00071.safetensors +3 -0
- model-00016-of-00071.safetensors +3 -0
- model-00017-of-00071.safetensors +3 -0
- model-00018-of-00071.safetensors +3 -0
- model-00019-of-00071.safetensors +3 -0
- model-00020-of-00071.safetensors +3 -0
- model-00021-of-00071.safetensors +3 -0
- model-00022-of-00071.safetensors +3 -0
- model-00023-of-00071.safetensors +3 -0
- model-00024-of-00071.safetensors +3 -0
- model-00025-of-00071.safetensors +3 -0
- model-00026-of-00071.safetensors +3 -0
- model-00027-of-00071.safetensors +3 -0
- model-00028-of-00071.safetensors +3 -0
- model-00029-of-00071.safetensors +3 -0
- model-00030-of-00071.safetensors +3 -0
- model-00031-of-00071.safetensors +3 -0
- model-00032-of-00071.safetensors +3 -0
- model-00033-of-00071.safetensors +3 -0
- model-00034-of-00071.safetensors +3 -0
- model-00035-of-00071.safetensors +3 -0
- model-00036-of-00071.safetensors +3 -0
- model-00037-of-00071.safetensors +3 -0
- model-00038-of-00071.safetensors +3 -0
- model-00039-of-00071.safetensors +3 -0
- model-00040-of-00071.safetensors +3 -0
- model-00041-of-00071.safetensors +3 -0
- model-00042-of-00071.safetensors +3 -0
- model-00043-of-00071.safetensors +3 -0
- model-00044-of-00071.safetensors +3 -0
- model-00045-of-00071.safetensors +3 -0
- model-00046-of-00071.safetensors +3 -0
- model-00047-of-00071.safetensors +3 -0
README.md
ADDED
@@ -0,0 +1,41 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
---
|
2 |
+
base_model:
|
3 |
+
- MidoriUnko/Behemoth-v2.2-Magnum-v4-123B
|
4 |
+
- MarsupialAI/Monstral-123B-v2
|
5 |
+
library_name: transformers
|
6 |
+
tags:
|
7 |
+
- mergekit
|
8 |
+
- merge
|
9 |
+
|
10 |
+
---
|
11 |
+
# merged_model
|
12 |
+
|
13 |
+
This is a merge of pre-trained language models created using [mergekit](https://github.com/cg123/mergekit).
|
14 |
+
|
15 |
+
## Merge Details
|
16 |
+
### Merge Method
|
17 |
+
|
18 |
+
This model was merged using the passthrough merge method.
|
19 |
+
|
20 |
+
### Models Merged
|
21 |
+
|
22 |
+
The following models were included in the merge:
|
23 |
+
* [MidoriUnko/Behemoth-v2.2-Magnum-v4-123B](https://huggingface.co/MidoriUnko/Behemoth-v2.2-Magnum-v4-123B)
|
24 |
+
* [MarsupialAI/Monstral-123B-v2](https://huggingface.co/MarsupialAI/Monstral-123B-v2)
|
25 |
+
|
26 |
+
### Configuration
|
27 |
+
|
28 |
+
The following YAML configuration was used to produce this model:
|
29 |
+
|
30 |
+
```yaml
|
31 |
+
slices:
|
32 |
+
- sources:
|
33 |
+
- model: MarsupialAI/Monstral-123B-v2
|
34 |
+
layer_range: [0, 61]
|
35 |
+
- sources:
|
36 |
+
- model: MidoriUnko/Behemoth-v2.2-Magnum-v4-123B
|
37 |
+
layer_range: [27, 88]
|
38 |
+
merge_method: passthrough
|
39 |
+
dtype: float16
|
40 |
+
name: monstral
|
41 |
+
```
|
config.json
ADDED
@@ -0,0 +1,27 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"_name_or_path": "MidoriUnko/Behemoth-v2.2-Magnum-v4-123B",
|
3 |
+
"architectures": [
|
4 |
+
"MistralForCausalLM"
|
5 |
+
],
|
6 |
+
"attention_dropout": 0.0,
|
7 |
+
"bos_token_id": 1,
|
8 |
+
"eos_token_id": 2,
|
9 |
+
"head_dim": 128,
|
10 |
+
"hidden_act": "silu",
|
11 |
+
"hidden_size": 12288,
|
12 |
+
"initializer_range": 0.02,
|
13 |
+
"intermediate_size": 28672,
|
14 |
+
"max_position_embeddings": 131072,
|
15 |
+
"model_type": "mistral",
|
16 |
+
"num_attention_heads": 96,
|
17 |
+
"num_hidden_layers": 122,
|
18 |
+
"num_key_value_heads": 8,
|
19 |
+
"rms_norm_eps": 1e-05,
|
20 |
+
"rope_theta": 1000000.0,
|
21 |
+
"sliding_window": null,
|
22 |
+
"tie_word_embeddings": false,
|
23 |
+
"torch_dtype": "float16",
|
24 |
+
"transformers_version": "4.47.1",
|
25 |
+
"use_cache": true,
|
26 |
+
"vocab_size": 32768
|
27 |
+
}
|
mergekit_config.yml
ADDED
@@ -0,0 +1,10 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
slices:
|
2 |
+
- sources:
|
3 |
+
- model: MarsupialAI/Monstral-123B-v2
|
4 |
+
layer_range: [0, 61]
|
5 |
+
- sources:
|
6 |
+
- model: MidoriUnko/Behemoth-v2.2-Magnum-v4-123B
|
7 |
+
layer_range: [27, 88]
|
8 |
+
merge_method: passthrough
|
9 |
+
dtype: float16
|
10 |
+
name: monstral
|
model-00001-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:41ff74477d0081835b213aec3d7ecc35a21e7e53467c58566048c5392602261d
|
3 |
+
size 4378928488
|
model-00002-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:2e17a9b2a1e4500c5a9fd7a219e880a8670bf688b68dc9ab29618a0aec74e83b
|
3 |
+
size 4907411072
|
model-00003-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:e1b6590c5b4e6d37323ce3d6666881e987e4a2321d3d9621820bb0319079defe
|
3 |
+
size 4806747888
|
model-00004-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:220eb35df35d5c24c31528ecbeb3570bc7ee464ccf4b1e50654e50dae2a719b9
|
3 |
+
size 4831938528
|
model-00005-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:063e5db6771b651c709190ef1d0ffef89dae2204648044ccaad30d258d23cc95
|
3 |
+
size 4831938536
|
model-00006-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:2d54039dd70f461b164c0905f9203f8aaf9bf0921da3ccb6a48003ec5da824f4
|
3 |
+
size 4907411080
|
model-00007-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:0752ad241695d36a76a973463ffecba679e19c621863283f029c22f55481b225
|
3 |
+
size 4806747888
|
model-00008-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:168640080680cb0f534d8e89ccf2e51e24526147c5f387fe1e52f5f43119e789
|
3 |
+
size 4831938520
|
model-00009-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:4c28a21aecb702c8fc049b8ae20d4f57672f09bd392cd58e815d4a0efeabc56f
|
3 |
+
size 4831938536
|
model-00010-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5291e0cabd1baaa6b7871a64207c367e633817c95bb3bffe8f965b72474d7d50
|
3 |
+
size 4907411080
|
model-00011-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:7e902aa69bafc0553661efd2fbb12e6a8564a4bc33b30567680ca39a38e26c2b
|
3 |
+
size 4806747888
|
model-00012-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:58cddae4428e5bf13d1c15bf39c976613ecaacb42c9a5f26a477e3fdd44f3be1
|
3 |
+
size 4831963216
|
model-00013-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:59955349d559bde405fce51dc53a35dfe47e68213255250f1e34395bee214c75
|
3 |
+
size 4831938536
|
model-00014-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:6d529a46d9340ae00c0c4a36d4ef392a94dbad0374a1bd8ba513a2e2f638c441
|
3 |
+
size 4882220448
|
model-00015-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b7e1c3a946cb71f7e7bf0e876841faa8faa373272a518866f89d8da24b8aadab
|
3 |
+
size 4932601704
|
model-00016-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f2ade0c25940aaadf1ca57a8f6a888e162fbfb71dcc49b09b1020d95e80d2f10
|
3 |
+
size 4731275336
|
model-00017-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:90d1f62e5519192c2c186940fc4af6754af749ea0c1941c6703f701e442a9c29
|
3 |
+
size 4831938536
|
model-00018-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c99d5a679a1534de4400b4bb6845f04d271a2add9b27839205dcfa8c55781a64
|
3 |
+
size 4882220448
|
model-00019-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:8d22ad9793ddd04711c57c8d42116a883b74cc8324a62a800f3afff0ced1154a
|
3 |
+
size 4932601704
|
model-00020-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:2434ef85595f843bfc6f185da916e5f2b83483815370181945554096561b6685
|
3 |
+
size 4781557256
|
model-00021-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:6bbc51e614472fbf0b9368bc02021394b3b59298eaf50a10feaf954c5254dc39
|
3 |
+
size 4831938536
|
model-00022-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f9f5016c6ed0c9e0ddf4292c6c927e2d64261a006501bf04f39366f3ccb617ea
|
3 |
+
size 4831938528
|
model-00023-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:e7af9048d9bab870396d90c085f2af2683259ca83b61323381d3226dcded70b5
|
3 |
+
size 4831938536
|
model-00024-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1d4af440d36d9235362219d523968edf54650d4e7bb384699c0be6db58f93de6
|
3 |
+
size 4831938536
|
model-00025-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:82c84d77e159c8b0d377fee44656b148ba16341082447289cb98745603b6ef97
|
3 |
+
size 4831938536
|
model-00026-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:de8eb365231249dde44260095d69769a85496aaa78541e53111c90341332d810
|
3 |
+
size 4882220448
|
model-00027-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c83d031ba86c00b5326e4b219e0e33b27e9cbc519abfdee678e704f1f74df4cc
|
3 |
+
size 4932601704
|
model-00028-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c8047dcd4899685d47a593d1c48c196a8e2af8108405b7f83991478dba53e72d
|
3 |
+
size 4731275336
|
model-00029-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:3dc036e5ba00c401285f0aaa699cf2ba548d8de59050d2bd39e75cb516becb58
|
3 |
+
size 4831938536
|
model-00030-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ab5aacf087ea755c13f12af95ff66dbe855dd6943dec437a31f1c872f37ba240
|
3 |
+
size 4882220448
|
model-00031-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:0c11670db4981aaf8deb379d8edacbc8287efb384831175c097ef52907cf0449
|
3 |
+
size 4932601704
|
model-00032-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:037e40855f176417cc9429187695832bd2747a98c9d42b25fd9bd36ab8e25821
|
3 |
+
size 4781557256
|
model-00033-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:0ca8314ce616244fd3a1c0856d7bb50829d89522c44e75efdca792d81b4205f8
|
3 |
+
size 4831938536
|
model-00034-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:df730ec37a364ea2ad7cbbec8ec5ba0bf61e9965be75932f0978cc5481e15e6f
|
3 |
+
size 4831938528
|
model-00035-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:0e010f1697cbe58eacb8ff4d6d968efdeb21da11dc10966c9f635b8f377418bc
|
3 |
+
size 4831938536
|
model-00036-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1508529063dac3c48db4ff8061f5f90c865bd5b2b31cc5a980a44d6aeeaaacfb
|
3 |
+
size 4831938536
|
model-00037-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:3961df4fb44f835581e922ef688ac5e46f27580b45f8d83b511057b0b6a5ae44
|
3 |
+
size 4831938536
|
model-00038-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:7e5c4941aaa5a8f1b2a2bcb154c99b1e7bd3a42307014a38bcfb80e5891058f9
|
3 |
+
size 4882220448
|
model-00039-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:58e516c98e891f316bec3a9600f8c011c0b3b0af0c5d351da57bbf5ea1ed645d
|
3 |
+
size 4932601704
|
model-00040-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a23ee3ce538e7af7b7b7394245364118dd4b44c9b52c1ef96d3d66c6a5420548
|
3 |
+
size 4731275336
|
model-00041-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:127a71e61b7664091c97f4599314beb8c7419ee2bf68b20e49d285d85a423570
|
3 |
+
size 4831938536
|
model-00042-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c73700ec1285bdf56632325d68eacebd85275fd7fb32853844ed3e1c942fa681
|
3 |
+
size 4882220448
|
model-00043-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:dc88155924be1f6d868d9e11b012c177d6483970439e0f6345d47c7338765157
|
3 |
+
size 4932601704
|
model-00044-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b92f4010ac2310400bd8db84cf32ee4d96015036ae9a633910abb950d4dcea74
|
3 |
+
size 4781557256
|
model-00045-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:60aaba8220c984f3082490b4fc88c1f7daf5145498d12a574ada779d61b3d3c2
|
3 |
+
size 4831938536
|
model-00046-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a6ef13c021190527262bff43170e2b88209742c9f111ae278e30b86230e015a0
|
3 |
+
size 4831938528
|
model-00047-of-00071.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:36b5eb72fff0a1d52d463fcf841b58aca747c8351ad8d3fcf94b709653592e65
|
3 |
+
size 4831938536
|