Shamane commited on
Commit
f007fdf
·
verified ·
1 Parent(s): 5bf35d0

Upload MistralForCausalLM

Browse files
config.json ADDED
@@ -0,0 +1,27 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "mistralai/Mistral-7B-v0.1",
3
+ "architectures": [
4
+ "MistralForCausalLM"
5
+ ],
6
+ "attention_dropout": 0.0,
7
+ "bos_token_id": 1,
8
+ "eos_token_id": 2,
9
+ "head_dim": 128,
10
+ "hidden_act": "silu",
11
+ "hidden_size": 4096,
12
+ "initializer_range": 0.02,
13
+ "intermediate_size": 14336,
14
+ "max_position_embeddings": 32768,
15
+ "model_type": "mistral",
16
+ "num_attention_heads": 32,
17
+ "num_hidden_layers": 32,
18
+ "num_key_value_heads": 8,
19
+ "rms_norm_eps": 1e-05,
20
+ "rope_theta": 10000.0,
21
+ "sliding_window": 4096,
22
+ "tie_word_embeddings": false,
23
+ "torch_dtype": "bfloat16",
24
+ "transformers_version": "4.44.2",
25
+ "use_cache": true,
26
+ "vocab_size": 32000
27
+ }
generation_config.json ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ {
2
+ "_from_model_config": true,
3
+ "bos_token_id": 1,
4
+ "eos_token_id": 2,
5
+ "transformers_version": "4.44.2"
6
+ }
model-00001-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5dec07ffbc2845afc7a03c12749010814d2ca30103607f4e9f5b2cd2a8636f88
3
+ size 4992325120
model-00002-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a6fbe486e1140f4249c8e7bd089d4ff042be3ad4be58cb82f16226da4fb5e8c7
3
+ size 4984151552
model-00003-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:eba6e65fa2a543acc32b82d1e9af15af4ddff8050a93d4d98c67866ad4d8340b
3
+ size 4916809056
model-00004-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:18d935c3e4fe2662ceebfc8003de6c5ae41748a190973174da6c7263e7502562
3
+ size 4883521784
model-00005-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5b485b81dcc382ed6c803ebf3905387ec88218418305a0fa9a0125040cb19a63
3
+ size 4883406784
model-00006-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:f77ac16faf284ced1b3cc35c89dc62542c653ec750261d72a2cab722b202c5d2
3
+ size 4883488680
model-00007-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:581329837bb02865062b47eea0b959052e597dddac1f3117676c7409e5afac49
3
+ size 4967374784
model-00008-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:03746d8b0f1e4b6a35cda48becc0f6d5682b190097576921d7fae177212c0c8e
3
+ size 4967306968
model-00009-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d323d403298290c92e5491a4b96a35e9b0a0b9f9332dbe888ab18547627d4569
3
+ size 4967241776
model-00010-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:b02ff5831e0328d76c9184cfe90911d80301797b56e9f6f2ddd26f7b5c6380e6
3
+ size 4883406776
model-00011-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:861f597d91e552e338842bbb1b42d116e98b357fe17568daced6096d13b1d92d
3
+ size 4883521896
model-00012-of-00012.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:61f52e2c730a206ed6eadaf35ad4d75a8090b2f9cc5e42eb211d0c55fb602d7f
3
+ size 3733730272
model.safetensors.index.json ADDED
The diff for this file is too large to render. See raw diff