at676 commited on
Commit
9085ca0
·
verified ·
1 Parent(s): bb9cc9b

Upload folder using huggingface_hub (#1)

Browse files

- b47a2ef875513374d9aa9b8815c619e55ed5f619944dcbcebc1166b393982a6c (aafdb6cb6c620324420260a9f39567dc24d1b6cc)
- a2171f9de7ff34c05cee394b4457f2a71881cdfd8ac6ae6731254420a4414fd3 (c44e3e6b3cbff19e8e32ababb695a8880f10474c)
- c87770b966e04019a8020cc149c1f52ef87d62ba447729d73b2e3c366a470167 (0270a3415434449f9ccb78f19f2b8ff43d3a351a)
- 456b62e6f0e669a5efff96d93f71e8df17da4c06fc59299a182ebcf1de4f1c61 (756a1b2dca45aea1b313b7f8ca8566881239972c)
- 7ed9fb67a03f0f9916003526844ab3401a3e4fd1616ff76d7ddaa76dfa97f56f (cde6df20a26b63ddd7a5d10d2e316807280b16fa)
- e8b3ea7d7637a82553547e6b075e29fb8db81237d8567482dfb413332d163948 (3b0bd6cf020364f21c67aff3722f5c2df5a9d574)
- c3d435085b5a67e486e5be722a81bfa9d292f68a59aa3646ca7f57b5eb46fdc0 (ea5e192b6a76cfa77cdf7ad680077adef81c76ef)
- f677188be3cb19fc24ef3db73af76cd716823b3bbcd47c581d19338700fad7f6 (f0c0780c171ccc5a5d2cd04a49a9dc939ac0f569)
- f6231d8dd71c01c5cd5ad795c7df126a15b62cf69b60086f092e578a03d0b119 (19b9aebd33b056370235c9f803f2358c8befe8dc)
- eb823efe0dfba450ffcbff1ce7dd94f8f2e6ca8a7b263804b201cb951e6a41e4 (181c4c9bfa93ec85461043b51a6d25a46a8688f1)
- 1ba4f14156c52b127b31ed60e003272b8f54fb71b3c4f273f5869fce30424827 (58c709fb6509345138f4ff5dc50cda317f83716b)
- b6f8bbdf9a05983059cab3f84c71d7eacb624b98388301b3cdff6b8e70edc517 (73950854678e81fcef4dacd30234f90271346b86)
- 8efb4530a2a03359bfc2da7ead3bf35c8c51e6cf0308d1a60e68c49f91046e4a (e7c861e2cdc66b595cbbeb419227de4c3c4bf2c1)
- ccc2ef3ea095287597bc9d0fd721a1480af2125891b8aef2d704c31834ea44c1 (2fef565859f06959e30d6d1142fa7fb261f5d119)
- 439750ab9e2e9fba3c939a2e8113cb8c675c97fceb6d3e61044ea3423729be59 (970a1081492f7c4096c08bd68fb686a6435cab2c)
- 1cdb3dee4ccd64261e72594d04b4393e546d34d73fa27c423c8183af3c9db69a (cacbfcdfeb5e05aa3511e038626da623deb272a6)
- 3082473893e67a6fc0cf7f9b4bc8b0ea40236efc9bbec3b42a54b5157484c7f5 (82fccbb4e7cdca160d260746b57bbb3aeb4908f1)
- c2593c2371a8a258d9698e43a9d69f031dbbd5f0b36ccea9282f3c256ab1e8e3 (8921485cf9be946ffe8d9d20c2f7743b2b6be9a8)
- 61c13666e916e0a7b3b6d568c11c764afd4cff8609bb73b32f9e540115cf4f3e (fa2a9da89312cfc1beaf8ef294309bf8446ca0bf)
- 474cf764366c65408f206141237a8ecc081974f62731b66b125ac7336b2e20ef (ca842afe475625f12ee5b4108f9ad536352b1bdc)
- e33ef69fd41e7bd3d9a4f4dea56b8e43ea069da720e1144827ce4417ef27e486 (205ebf9e8f586327254c80a4b3cd0b40194b3b82)
- b072e888c8061bc56b3b0c1fcea6c6f7b81a1980278d061bcc73fc193abecb88 (ddd06163722df024224e17d9e27fb2a6fa316658)
- 9ffdfb07c0411dd38e372b3ab60756f02eecef89a934e696962c3e4846639a6d (fec14fa04427a57e901eb94dbcde1d1209378b85)
- 157445ec43888980321172680ee2215167e46d7ceb34df765911b124a5f38730 (c37474cce555fe60ded7da1ea254ef19da13bcd1)

config.json ADDED
@@ -0,0 +1,51 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "meta-llama/Meta-Llama-3.1-405B-Instruct",
3
+ "architectures": [
4
+ "LlamaForCausalLM"
5
+ ],
6
+ "attention_bias": false,
7
+ "attention_dropout": 0.0,
8
+ "bos_token_id": 128000,
9
+ "eos_token_id": [
10
+ 128001,
11
+ 128008,
12
+ 128009
13
+ ],
14
+ "head_dim": 128,
15
+ "hidden_act": "silu",
16
+ "hidden_size": 16384,
17
+ "initializer_range": 0.02,
18
+ "intermediate_size": 53248,
19
+ "max_position_embeddings": 131072,
20
+ "mlp_bias": false,
21
+ "model_type": "llama",
22
+ "num_attention_heads": 128,
23
+ "num_hidden_layers": 126,
24
+ "num_key_value_heads": 8,
25
+ "pretraining_tp": 1,
26
+ "quip_params": {
27
+ "K": 2,
28
+ "L": 16,
29
+ "V": 2,
30
+ "codebook": "bitshift",
31
+ "codebook_version": 0,
32
+ "decode_mode": "quantlut_sym",
33
+ "td_x": 16,
34
+ "td_y": 16,
35
+ "tlut_bits": 9
36
+ },
37
+ "rms_norm_eps": 1e-05,
38
+ "rope_scaling": {
39
+ "factor": 8.0,
40
+ "high_freq_factor": 4.0,
41
+ "low_freq_factor": 1.0,
42
+ "original_max_position_embeddings": 8192,
43
+ "rope_type": "llama3"
44
+ },
45
+ "rope_theta": 500000.0,
46
+ "tie_word_embeddings": false,
47
+ "torch_dtype": "bfloat16",
48
+ "transformers_version": "4.45.2",
49
+ "use_cache": true,
50
+ "vocab_size": 128256
51
+ }
generation_config.json ADDED
@@ -0,0 +1,12 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token_id": 128000,
3
+ "do_sample": true,
4
+ "eos_token_id": [
5
+ 128001,
6
+ 128008,
7
+ 128009
8
+ ],
9
+ "temperature": 0.6,
10
+ "top_p": 0.9,
11
+ "transformers_version": "4.45.2"
12
+ }
model-00001-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:04d2984f2da1a7c01c0595613b447817dcd344988b2cd6ac0f57d6126327aa4f
3
+ size 4782285584
model-00002-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d51f0c2199bc5cf23697db0a238648cef87d6d17c8b34df1cff6b33b3d2b88dd
3
+ size 4787608080
model-00003-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:94a25ecdac67741797eee196b011abe7709ff4bf79777c543dced29f510bd7b2
3
+ size 4787608168
model-00004-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:96398cdec1da75e421dcfe46fdc3d0c087b73e6d9a8571ec9a3a04385019f107
3
+ size 4787608264
model-00005-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5f0ee5db899ba3624ba016cb2f9c82bffc79e60904157887bbd2fc2f7316d682
3
+ size 4787608264
model-00006-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:f2503f55df6d5423acbc71fa68c75e02cfc11acbdc59209d87fb2d9c54e51098
3
+ size 4787608264
model-00007-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:eefd91645f0568f1402ebaaaf7201d69c74228a945e81f6f8fbdc268a835ccec
3
+ size 4787608264
model-00008-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e121b7ffcba1c643a7732be1b96f949998110ad331d897aff90abe13dca1f512
3
+ size 4787608264
model-00009-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cfbb0a5bfd7bcdcceeb5309f5f48124855ca2d5c19bb61cd3e8cd0eaf0e56ef9
3
+ size 4787608264
model-00010-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:1419ec19a8f3b655f5c6cd0360469cffce28170ce349fc70ece55564600f02fd
3
+ size 4787608264
model-00011-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:42b742883f04ec0e41c7cd66a80279afd4469b55f8836e629d8b07764bffed1a
3
+ size 4787608264
model-00012-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:456ae5d512826235cb7d3a062afedc6e28c6855aa446b49c4e6361189692d738
3
+ size 4787608264
model-00013-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ca63f479ea15b39535657fc8b2bc066e5fb1b79c01b19a80aabd8b146e252d71
3
+ size 4787608264
model-00014-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c89464f248a54d52ffcce955ba39633cdca357181de938e1efe896ece9161306
3
+ size 4787608264
model-00015-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:44be9e9633e2d56c3f8c7b6ef175b02db7f744fe600900bd23b9d979cd38c904
3
+ size 4787608264
model-00016-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:8603ae05a449d5ffc10cd23612f1b65faf9f213a3aec960050f4ef2862644ec4
3
+ size 4787608264
model-00017-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e13ff34432b74effc7dade394f7d9e7ecf0130a7bdfe4a3a453bd2fb92ecb0b2
3
+ size 4787608264
model-00018-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:8f62f217e5eb9be257c394998c3d78c49f4d2daba10a3e657a2a5a47dd4e5b85
3
+ size 4787608344
model-00019-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:10e4129eda4f782a7d90675420ab83e38e5db872061581a869775f1481ecb68e
3
+ size 4787608440
model-00020-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7b2862a305cdb846160545e2524d0bbf5937e986a450ab37c1874a267349ccc9
3
+ size 4787608440
model-00021-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:9b661d5d847b890b6bf6593e0b8aad3e228aacde9ad886da48353509dc3494a7
3
+ size 4787608440
model-00022-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a0095d467f0be13a4934165c27ecc39fb0fe37aa2e1ee267b3e4ffa7d77c4a61
3
+ size 4208048472
model-00023-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:71b079a02069ebe5a5f4f8135815f136afb1d6fc4423620dcb2a14fa399cdf12
3
+ size 4202692736
model.safetensors.index.json ADDED
The diff for this file is too large to render. See raw diff