at676 commited on
Commit
93b6367
1 Parent(s): a14ac2e

Upload folder using huggingface_hub (#1)

Browse files

- f937bcdbef3a34a214ab3f1972260f1d54bf14b5cd8b0df79ab544f08f64659c (f8713410cd835a2b378c54378f6603a6ab0661db)
- 0f1f59b307ee71d20492da0c464b7eecdd959b01ddc6e5d8de90e1d90d52923b (b7d820d471b96f50b5258fec9bdec431a17ffb1e)
- efdfa6f4716bb89a5cecf2cdfd8e1020f568c47faa88ea60d02b1fbe2caf7ad2 (fb6d4bc9ef7e1762e4fa751372d8d03bdb7217eb)
- 81f04f58755195190e8c98d2a4b6a528e5356a516a0f7ac2afb6aeace2d4fa2e (069bca08c8aaf74db23885ccc6a73bc8d5992f2f)
- b852a895516baf4d1cf8098f6a7e79ce7f002df1ceb7f90bd55d7334916c13d3 (e58d6ad2e07a91680181d4e23c04a11adbe8227c)
- a89cee75c3cdc20b548d2e2578862d6042c1e0369f20902357f5bdb13071a62b (d96f173a25345cc5c546046dc44397d687339024)
- 6b680dff147fa83ed41e0198ed049b2295b372ec2b94279d128843ddc5783809 (749b82c08446ea0374eb4b590136446410565acb)
- b15293ac1a5fd6b412f4e2b0558bb2d1d1e51f803fa19f3056944e893e96c2a1 (17f89ea695b9fad77a4c4d93b6c0964097e67412)

config.json ADDED
@@ -0,0 +1,40 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "meta-llama/Llama-2-70b-hf",
3
+ "architectures": [
4
+ "LlamaForCausalLM"
5
+ ],
6
+ "attention_bias": false,
7
+ "attention_dropout": 0.0,
8
+ "bos_token_id": 1,
9
+ "eos_token_id": 2,
10
+ "hidden_act": "silu",
11
+ "hidden_size": 8192,
12
+ "initializer_range": 0.02,
13
+ "intermediate_size": 28672,
14
+ "max_position_embeddings": 4096,
15
+ "mlp_bias": false,
16
+ "model_type": "llama",
17
+ "num_attention_heads": 64,
18
+ "num_hidden_layers": 80,
19
+ "num_key_value_heads": 8,
20
+ "pretraining_tp": 1,
21
+ "quip_params": {
22
+ "K": 4,
23
+ "L": 16,
24
+ "V": 2,
25
+ "codebook": "bitshift",
26
+ "codebook_version": 0,
27
+ "decode_mode": "quantlut_sym",
28
+ "td_x": 16,
29
+ "td_y": 16,
30
+ "tlut_bits": 9
31
+ },
32
+ "rms_norm_eps": 1e-05,
33
+ "rope_scaling": null,
34
+ "rope_theta": 10000.0,
35
+ "tie_word_embeddings": false,
36
+ "torch_dtype": "float16",
37
+ "transformers_version": "4.44.2",
38
+ "use_cache": true,
39
+ "vocab_size": 32000
40
+ }
generation_config.json ADDED
@@ -0,0 +1,10 @@
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token_id": 1,
3
+ "do_sample": true,
4
+ "eos_token_id": 2,
5
+ "max_length": 4096,
6
+ "pad_token_id": 0,
7
+ "temperature": 0.6,
8
+ "top_p": 0.9,
9
+ "transformers_version": "4.44.2"
10
+ }
model-00001-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ff641d04b993e03d72cfda0cd02bd3efdeef9e4cadadf24c91d0f8e69981e125
3
+ size 4883545208
model-00002-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c1893a83027eaec9ed6bfd2b71fa084514d11200179179b6a46f32252dfb7998
3
+ size 4947116832
model-00003-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cfa0cbb45ef900a4463240ecfae840e5a00e150a742418d6d7a28b8d12202dc3
3
+ size 4905181440
model-00004-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0ac1c667027495671ed5ef78a722b8067a9c0a19d0dd14cd0c7a2280cf992c1c
3
+ size 4947116832
model-00005-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:09ca3c164f58d1f785eb22214d77f1e096a6c0c8fe3432b69f65401fafbab76e
3
+ size 4905181440
model-00006-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:6286a2c7b6cd0a43b30520ceb1c12b43a9ecb10b68714c63a16b8be2ea98b026
3
+ size 4947116832
model-00007-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ba28bda1b48489f7a1b6ce30ebe757b5e17f7be2c80f6dc25047fe1a527dc998
3
+ size 4905181440
model-00008-of-00008.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5bfd1ac68371f3589499a3ad0e8a5db8f554690f9e153732ac3e87ef75510a2d
3
+ size 877016624
model.safetensors.index.json ADDED
The diff for this file is too large to render. See raw diff