Upload folder using huggingface_hub (#1)
Browse files- f937bcdbef3a34a214ab3f1972260f1d54bf14b5cd8b0df79ab544f08f64659c (f8713410cd835a2b378c54378f6603a6ab0661db)
- 0f1f59b307ee71d20492da0c464b7eecdd959b01ddc6e5d8de90e1d90d52923b (b7d820d471b96f50b5258fec9bdec431a17ffb1e)
- efdfa6f4716bb89a5cecf2cdfd8e1020f568c47faa88ea60d02b1fbe2caf7ad2 (fb6d4bc9ef7e1762e4fa751372d8d03bdb7217eb)
- 81f04f58755195190e8c98d2a4b6a528e5356a516a0f7ac2afb6aeace2d4fa2e (069bca08c8aaf74db23885ccc6a73bc8d5992f2f)
- b852a895516baf4d1cf8098f6a7e79ce7f002df1ceb7f90bd55d7334916c13d3 (e58d6ad2e07a91680181d4e23c04a11adbe8227c)
- a89cee75c3cdc20b548d2e2578862d6042c1e0369f20902357f5bdb13071a62b (d96f173a25345cc5c546046dc44397d687339024)
- 6b680dff147fa83ed41e0198ed049b2295b372ec2b94279d128843ddc5783809 (749b82c08446ea0374eb4b590136446410565acb)
- b15293ac1a5fd6b412f4e2b0558bb2d1d1e51f803fa19f3056944e893e96c2a1 (17f89ea695b9fad77a4c4d93b6c0964097e67412)
- config.json +40 -0
- generation_config.json +10 -0
- model-00001-of-00008.safetensors +3 -0
- model-00002-of-00008.safetensors +3 -0
- model-00003-of-00008.safetensors +3 -0
- model-00004-of-00008.safetensors +3 -0
- model-00005-of-00008.safetensors +3 -0
- model-00006-of-00008.safetensors +3 -0
- model-00007-of-00008.safetensors +3 -0
- model-00008-of-00008.safetensors +3 -0
- model.safetensors.index.json +0 -0
@@ -0,0 +1,40 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"_name_or_path": "meta-llama/Llama-2-70b-hf",
|
3 |
+
"architectures": [
|
4 |
+
"LlamaForCausalLM"
|
5 |
+
],
|
6 |
+
"attention_bias": false,
|
7 |
+
"attention_dropout": 0.0,
|
8 |
+
"bos_token_id": 1,
|
9 |
+
"eos_token_id": 2,
|
10 |
+
"hidden_act": "silu",
|
11 |
+
"hidden_size": 8192,
|
12 |
+
"initializer_range": 0.02,
|
13 |
+
"intermediate_size": 28672,
|
14 |
+
"max_position_embeddings": 4096,
|
15 |
+
"mlp_bias": false,
|
16 |
+
"model_type": "llama",
|
17 |
+
"num_attention_heads": 64,
|
18 |
+
"num_hidden_layers": 80,
|
19 |
+
"num_key_value_heads": 8,
|
20 |
+
"pretraining_tp": 1,
|
21 |
+
"quip_params": {
|
22 |
+
"K": 4,
|
23 |
+
"L": 16,
|
24 |
+
"V": 2,
|
25 |
+
"codebook": "bitshift",
|
26 |
+
"codebook_version": 0,
|
27 |
+
"decode_mode": "quantlut_sym",
|
28 |
+
"td_x": 16,
|
29 |
+
"td_y": 16,
|
30 |
+
"tlut_bits": 9
|
31 |
+
},
|
32 |
+
"rms_norm_eps": 1e-05,
|
33 |
+
"rope_scaling": null,
|
34 |
+
"rope_theta": 10000.0,
|
35 |
+
"tie_word_embeddings": false,
|
36 |
+
"torch_dtype": "float16",
|
37 |
+
"transformers_version": "4.44.2",
|
38 |
+
"use_cache": true,
|
39 |
+
"vocab_size": 32000
|
40 |
+
}
|
@@ -0,0 +1,10 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"bos_token_id": 1,
|
3 |
+
"do_sample": true,
|
4 |
+
"eos_token_id": 2,
|
5 |
+
"max_length": 4096,
|
6 |
+
"pad_token_id": 0,
|
7 |
+
"temperature": 0.6,
|
8 |
+
"top_p": 0.9,
|
9 |
+
"transformers_version": "4.44.2"
|
10 |
+
}
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ff641d04b993e03d72cfda0cd02bd3efdeef9e4cadadf24c91d0f8e69981e125
|
3 |
+
size 4883545208
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c1893a83027eaec9ed6bfd2b71fa084514d11200179179b6a46f32252dfb7998
|
3 |
+
size 4947116832
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:cfa0cbb45ef900a4463240ecfae840e5a00e150a742418d6d7a28b8d12202dc3
|
3 |
+
size 4905181440
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:0ac1c667027495671ed5ef78a722b8067a9c0a19d0dd14cd0c7a2280cf992c1c
|
3 |
+
size 4947116832
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:09ca3c164f58d1f785eb22214d77f1e096a6c0c8fe3432b69f65401fafbab76e
|
3 |
+
size 4905181440
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:6286a2c7b6cd0a43b30520ceb1c12b43a9ecb10b68714c63a16b8be2ea98b026
|
3 |
+
size 4947116832
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ba28bda1b48489f7a1b6ce30ebe757b5e17f7be2c80f6dc25047fe1a527dc998
|
3 |
+
size 4905181440
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5bfd1ac68371f3589499a3ad0e8a5db8f554690f9e153732ac3e87ef75510a2d
|
3 |
+
size 877016624
|
The diff for this file is too large to render.
See raw diff
|
|