Upload folder using huggingface_hub (#1)
Browse files- b47a2ef875513374d9aa9b8815c619e55ed5f619944dcbcebc1166b393982a6c (aafdb6cb6c620324420260a9f39567dc24d1b6cc)
- a2171f9de7ff34c05cee394b4457f2a71881cdfd8ac6ae6731254420a4414fd3 (c44e3e6b3cbff19e8e32ababb695a8880f10474c)
- c87770b966e04019a8020cc149c1f52ef87d62ba447729d73b2e3c366a470167 (0270a3415434449f9ccb78f19f2b8ff43d3a351a)
- 456b62e6f0e669a5efff96d93f71e8df17da4c06fc59299a182ebcf1de4f1c61 (756a1b2dca45aea1b313b7f8ca8566881239972c)
- 7ed9fb67a03f0f9916003526844ab3401a3e4fd1616ff76d7ddaa76dfa97f56f (cde6df20a26b63ddd7a5d10d2e316807280b16fa)
- e8b3ea7d7637a82553547e6b075e29fb8db81237d8567482dfb413332d163948 (3b0bd6cf020364f21c67aff3722f5c2df5a9d574)
- c3d435085b5a67e486e5be722a81bfa9d292f68a59aa3646ca7f57b5eb46fdc0 (ea5e192b6a76cfa77cdf7ad680077adef81c76ef)
- f677188be3cb19fc24ef3db73af76cd716823b3bbcd47c581d19338700fad7f6 (f0c0780c171ccc5a5d2cd04a49a9dc939ac0f569)
- f6231d8dd71c01c5cd5ad795c7df126a15b62cf69b60086f092e578a03d0b119 (19b9aebd33b056370235c9f803f2358c8befe8dc)
- eb823efe0dfba450ffcbff1ce7dd94f8f2e6ca8a7b263804b201cb951e6a41e4 (181c4c9bfa93ec85461043b51a6d25a46a8688f1)
- 1ba4f14156c52b127b31ed60e003272b8f54fb71b3c4f273f5869fce30424827 (58c709fb6509345138f4ff5dc50cda317f83716b)
- b6f8bbdf9a05983059cab3f84c71d7eacb624b98388301b3cdff6b8e70edc517 (73950854678e81fcef4dacd30234f90271346b86)
- 8efb4530a2a03359bfc2da7ead3bf35c8c51e6cf0308d1a60e68c49f91046e4a (e7c861e2cdc66b595cbbeb419227de4c3c4bf2c1)
- ccc2ef3ea095287597bc9d0fd721a1480af2125891b8aef2d704c31834ea44c1 (2fef565859f06959e30d6d1142fa7fb261f5d119)
- 439750ab9e2e9fba3c939a2e8113cb8c675c97fceb6d3e61044ea3423729be59 (970a1081492f7c4096c08bd68fb686a6435cab2c)
- 1cdb3dee4ccd64261e72594d04b4393e546d34d73fa27c423c8183af3c9db69a (cacbfcdfeb5e05aa3511e038626da623deb272a6)
- 3082473893e67a6fc0cf7f9b4bc8b0ea40236efc9bbec3b42a54b5157484c7f5 (82fccbb4e7cdca160d260746b57bbb3aeb4908f1)
- c2593c2371a8a258d9698e43a9d69f031dbbd5f0b36ccea9282f3c256ab1e8e3 (8921485cf9be946ffe8d9d20c2f7743b2b6be9a8)
- 61c13666e916e0a7b3b6d568c11c764afd4cff8609bb73b32f9e540115cf4f3e (fa2a9da89312cfc1beaf8ef294309bf8446ca0bf)
- 474cf764366c65408f206141237a8ecc081974f62731b66b125ac7336b2e20ef (ca842afe475625f12ee5b4108f9ad536352b1bdc)
- e33ef69fd41e7bd3d9a4f4dea56b8e43ea069da720e1144827ce4417ef27e486 (205ebf9e8f586327254c80a4b3cd0b40194b3b82)
- b072e888c8061bc56b3b0c1fcea6c6f7b81a1980278d061bcc73fc193abecb88 (ddd06163722df024224e17d9e27fb2a6fa316658)
- 9ffdfb07c0411dd38e372b3ab60756f02eecef89a934e696962c3e4846639a6d (fec14fa04427a57e901eb94dbcde1d1209378b85)
- 157445ec43888980321172680ee2215167e46d7ceb34df765911b124a5f38730 (c37474cce555fe60ded7da1ea254ef19da13bcd1)
- config.json +51 -0
- generation_config.json +12 -0
- model-00001-of-00023.safetensors +3 -0
- model-00002-of-00023.safetensors +3 -0
- model-00003-of-00023.safetensors +3 -0
- model-00004-of-00023.safetensors +3 -0
- model-00005-of-00023.safetensors +3 -0
- model-00006-of-00023.safetensors +3 -0
- model-00007-of-00023.safetensors +3 -0
- model-00008-of-00023.safetensors +3 -0
- model-00009-of-00023.safetensors +3 -0
- model-00010-of-00023.safetensors +3 -0
- model-00011-of-00023.safetensors +3 -0
- model-00012-of-00023.safetensors +3 -0
- model-00013-of-00023.safetensors +3 -0
- model-00014-of-00023.safetensors +3 -0
- model-00015-of-00023.safetensors +3 -0
- model-00016-of-00023.safetensors +3 -0
- model-00017-of-00023.safetensors +3 -0
- model-00018-of-00023.safetensors +3 -0
- model-00019-of-00023.safetensors +3 -0
- model-00020-of-00023.safetensors +3 -0
- model-00021-of-00023.safetensors +3 -0
- model-00022-of-00023.safetensors +3 -0
- model-00023-of-00023.safetensors +3 -0
- model.safetensors.index.json +0 -0
@@ -0,0 +1,51 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"_name_or_path": "meta-llama/Meta-Llama-3.1-405B-Instruct",
|
3 |
+
"architectures": [
|
4 |
+
"LlamaForCausalLM"
|
5 |
+
],
|
6 |
+
"attention_bias": false,
|
7 |
+
"attention_dropout": 0.0,
|
8 |
+
"bos_token_id": 128000,
|
9 |
+
"eos_token_id": [
|
10 |
+
128001,
|
11 |
+
128008,
|
12 |
+
128009
|
13 |
+
],
|
14 |
+
"head_dim": 128,
|
15 |
+
"hidden_act": "silu",
|
16 |
+
"hidden_size": 16384,
|
17 |
+
"initializer_range": 0.02,
|
18 |
+
"intermediate_size": 53248,
|
19 |
+
"max_position_embeddings": 131072,
|
20 |
+
"mlp_bias": false,
|
21 |
+
"model_type": "llama",
|
22 |
+
"num_attention_heads": 128,
|
23 |
+
"num_hidden_layers": 126,
|
24 |
+
"num_key_value_heads": 8,
|
25 |
+
"pretraining_tp": 1,
|
26 |
+
"quip_params": {
|
27 |
+
"K": 2,
|
28 |
+
"L": 16,
|
29 |
+
"V": 2,
|
30 |
+
"codebook": "bitshift",
|
31 |
+
"codebook_version": 0,
|
32 |
+
"decode_mode": "quantlut_sym",
|
33 |
+
"td_x": 16,
|
34 |
+
"td_y": 16,
|
35 |
+
"tlut_bits": 9
|
36 |
+
},
|
37 |
+
"rms_norm_eps": 1e-05,
|
38 |
+
"rope_scaling": {
|
39 |
+
"factor": 8.0,
|
40 |
+
"high_freq_factor": 4.0,
|
41 |
+
"low_freq_factor": 1.0,
|
42 |
+
"original_max_position_embeddings": 8192,
|
43 |
+
"rope_type": "llama3"
|
44 |
+
},
|
45 |
+
"rope_theta": 500000.0,
|
46 |
+
"tie_word_embeddings": false,
|
47 |
+
"torch_dtype": "bfloat16",
|
48 |
+
"transformers_version": "4.45.2",
|
49 |
+
"use_cache": true,
|
50 |
+
"vocab_size": 128256
|
51 |
+
}
|
@@ -0,0 +1,12 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"bos_token_id": 128000,
|
3 |
+
"do_sample": true,
|
4 |
+
"eos_token_id": [
|
5 |
+
128001,
|
6 |
+
128008,
|
7 |
+
128009
|
8 |
+
],
|
9 |
+
"temperature": 0.6,
|
10 |
+
"top_p": 0.9,
|
11 |
+
"transformers_version": "4.45.2"
|
12 |
+
}
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:04d2984f2da1a7c01c0595613b447817dcd344988b2cd6ac0f57d6126327aa4f
|
3 |
+
size 4782285584
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:d51f0c2199bc5cf23697db0a238648cef87d6d17c8b34df1cff6b33b3d2b88dd
|
3 |
+
size 4787608080
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:94a25ecdac67741797eee196b011abe7709ff4bf79777c543dced29f510bd7b2
|
3 |
+
size 4787608168
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:96398cdec1da75e421dcfe46fdc3d0c087b73e6d9a8571ec9a3a04385019f107
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5f0ee5db899ba3624ba016cb2f9c82bffc79e60904157887bbd2fc2f7316d682
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f2503f55df6d5423acbc71fa68c75e02cfc11acbdc59209d87fb2d9c54e51098
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:eefd91645f0568f1402ebaaaf7201d69c74228a945e81f6f8fbdc268a835ccec
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:e121b7ffcba1c643a7732be1b96f949998110ad331d897aff90abe13dca1f512
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:cfbb0a5bfd7bcdcceeb5309f5f48124855ca2d5c19bb61cd3e8cd0eaf0e56ef9
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1419ec19a8f3b655f5c6cd0360469cffce28170ce349fc70ece55564600f02fd
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:42b742883f04ec0e41c7cd66a80279afd4469b55f8836e629d8b07764bffed1a
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:456ae5d512826235cb7d3a062afedc6e28c6855aa446b49c4e6361189692d738
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ca63f479ea15b39535657fc8b2bc066e5fb1b79c01b19a80aabd8b146e252d71
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c89464f248a54d52ffcce955ba39633cdca357181de938e1efe896ece9161306
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:44be9e9633e2d56c3f8c7b6ef175b02db7f744fe600900bd23b9d979cd38c904
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:8603ae05a449d5ffc10cd23612f1b65faf9f213a3aec960050f4ef2862644ec4
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:e13ff34432b74effc7dade394f7d9e7ecf0130a7bdfe4a3a453bd2fb92ecb0b2
|
3 |
+
size 4787608264
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:8f62f217e5eb9be257c394998c3d78c49f4d2daba10a3e657a2a5a47dd4e5b85
|
3 |
+
size 4787608344
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:10e4129eda4f782a7d90675420ab83e38e5db872061581a869775f1481ecb68e
|
3 |
+
size 4787608440
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:7b2862a305cdb846160545e2524d0bbf5937e986a450ab37c1874a267349ccc9
|
3 |
+
size 4787608440
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:9b661d5d847b890b6bf6593e0b8aad3e228aacde9ad886da48353509dc3494a7
|
3 |
+
size 4787608440
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a0095d467f0be13a4934165c27ecc39fb0fe37aa2e1ee267b3e4ffa7d77c4a61
|
3 |
+
size 4208048472
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:71b079a02069ebe5a5f4f8135815f136afb1d6fc4423620dcb2a14fa399cdf12
|
3 |
+
size 4202692736
|
The diff for this file is too large to render.
See raw diff
|
|