at676 commited on
Commit
2ef22fb
·
verified ·
1 Parent(s): 6db864b

Upload folder using huggingface_hub (#2)

Browse files

- 76fbedef0bdfe37c6558b1c39a2a98f57d4ed352b2ef3e6648e5db1397f65c68 (58ec909e5afb76e278272e525941055e8bdf6d12)
- bdebc542a83866143c3565e82a97ab72cc5adb286ecbae4d821773797e7bb62e (a5daed980f9699dd43ec9b0945ead1915a89efbd)
- a586678e5021a3440855a0273dc2652354a9d6880ba30a4606f8b3d934b74cb3 (570ed53959e15da48500f35289f4f96c137b7ef3)
- 199248a46371eb41b2fb6dafc0e1a014f5dde5492e061e1c5eb3cb8846bf4760 (e0480cec45a51fe2586aefa0164f2b7635c83c2b)
- a2d14c42274f6a3855dbb4ff3c2092e5faef445beb75eb4201e8721c7054a2bf (e70ffc0a662dc7da9200f554be57804663ce673b)
- 0fd741c023a857cf6b7a5d7daada9b6b3ac01a74dbc701b5111de72c181298f3 (5fbac124a64b19c83914b2762b16f398bcabef43)
- 58ea601209371414158e15923c04ba601b5d60a671b8662f9130b0d4f1a16302 (d50df2cfa47a2b2ce9a4434659a925dae25653aa)
- 8322986c46f670f5d613499110fb0005f6122678ff7a57f7ebed73200f89f946 (44faadadf55231f983c0f34787d59766071db986)
- c3fc70a939aa79575db4bbfdbc6df8be78ffb1941522502b12305fe90d46539d (a576a6b37211efad779caa374e3ec117afb1c6dc)
- 9d42f567b2a97a864d177ffee8c9d962de47b26e4d23bf2ff11040f498f0b7c3 (3efe9b22811bf3078e2f5546954b0f878d86a695)
- 67c8b28633069b6fba0f3d7eeb0e06f4910b2daefe663eda129e76c74afe15f1 (165e64f3d4d1538dc36f0d6355c4aa0625653414)
- f6d126b3be03c87952b7e332b272eeda4cad535607126da9faefad8756350ef7 (3ccb8487eadc27cd96b91cf1f82d61b19cc3a06e)
- 6923d35211bc4e72afae711d500ed768db5c55d79f9a07923a1165d2ab11d85b (2520cb69ca4c3671d0aacf607d258a093f28e7ed)
- 882c83662637cd44efc425b3fe0e6388a04e365b3904101960f330197255277f (6f3a62561958b0de3399c759e8820e9587512ba6)
- e7dfa515edc3bb89b6dc880b657270c4f852d2db28c8c8e1de92c460affcfc59 (f75a7aa75b516e8b96a6db36ce8358f4ca09e76e)
- 6759d913d4b3f9075067862214e3ec617d9da4f6e7eb6d695fa3f46d4dff4885 (e01177feca6f1d627a532b7c1b12b054cc7e99c0)
- ab7ec00554a598585531249b3b432bf168c0f4b70fdbe8becdbc237f11ecd7ba (d3849c651b9c1ed80230b9256bf9884e0f93e5bd)
- 5bcfbd71d57a9a3dfe89f6101aa6f98271fdf659804f8dff1cc04a1bb6899363 (c2cf062e3a4d969df12676a4417056c3e556a652)
- 85139f3a354af8ea64b0864cd4adfd96102e754439be7ba02ae1a842326335cf (250bfa8fdf2723dd1eb99b777abdb33b6a787c2f)
- 5b7796430443a3abc8e738c521025696586e3161c440e4573f3b879ac64f167f (db9972a96b70c732f3cdea13373c3dd75493b685)
- a4a89716b21c8e64c43daf91cb2cfa96a34de01048d6d68af4ca56ab13fdbd8d (8fd0d02c46bc51d5a1f3801e4a4fcef276ab80be)
- 5d2069cdd999f038fda5dcd3e0376027e92d2e16203e57bb757f695b40183769 (d60fdc506541b8e34584b13c81c0a82439ae6ba9)
- 9ffdfb07c0411dd38e372b3ab60756f02eecef89a934e696962c3e4846639a6d (5b930342c4054262db636eb017403f75fa6b1385)
- 3d89d5280c27b35551a5aaf12a296e4421066ef709e5dfb5f243f336d4d374dc (853066a7c80eeebd99262ec571f46f501e9e262b)

config.json ADDED
@@ -0,0 +1,52 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "meta-llama/Meta-Llama-3.1-405B-Instruct",
3
+ "architectures": [
4
+ "LlamaForCausalLM"
5
+ ],
6
+ "attention_bias": false,
7
+ "attention_dropout": 0.0,
8
+ "bos_token_id": 128000,
9
+ "eos_token_id": [
10
+ 128001,
11
+ 128008,
12
+ 128009
13
+ ],
14
+ "head_dim": 128,
15
+ "hidden_act": "silu",
16
+ "hidden_size": 16384,
17
+ "initializer_range": 0.02,
18
+ "intermediate_size": 53248,
19
+ "max_position_embeddings": 131072,
20
+ "mlp_bias": false,
21
+ "model_type": "llama",
22
+ "num_attention_heads": 128,
23
+ "num_hidden_layers": 126,
24
+ "num_key_value_heads": 8,
25
+ "pretraining_tp": 1,
26
+ "quip_params": {
27
+ "K": 2,
28
+ "L": 16,
29
+ "V": 2,
30
+ "codebook": "bitshift",
31
+ "codebook_version": 0,
32
+ "decode_mode": "quantlut_sym",
33
+ "split_for_tp": true,
34
+ "td_x": 16,
35
+ "td_y": 16,
36
+ "tlut_bits": 9
37
+ },
38
+ "rms_norm_eps": 1e-05,
39
+ "rope_scaling": {
40
+ "factor": 8.0,
41
+ "high_freq_factor": 4.0,
42
+ "low_freq_factor": 1.0,
43
+ "original_max_position_embeddings": 8192,
44
+ "rope_type": "llama3"
45
+ },
46
+ "rope_theta": 500000.0,
47
+ "tie_word_embeddings": false,
48
+ "torch_dtype": "bfloat16",
49
+ "transformers_version": "4.45.2",
50
+ "use_cache": true,
51
+ "vocab_size": 128256
52
+ }
generation_config.json ADDED
@@ -0,0 +1,12 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token_id": 128000,
3
+ "do_sample": true,
4
+ "eos_token_id": [
5
+ 128001,
6
+ 128008,
7
+ 128009
8
+ ],
9
+ "temperature": 0.6,
10
+ "top_p": 0.9,
11
+ "transformers_version": "4.45.2"
12
+ }
model-00001-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:71ecc4b88297075b087cffc38786bd865496e8d94280ba52f771900e44c59be3
3
+ size 4782286936
model-00002-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e1e23c3cea3d9dffc29f5bd83b43f1d64745bd7d0a7589ac5ecb46293a609b8d
3
+ size 4787616320
model-00003-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a4b130b89d32aeaa5e3014f1a73da3728ba52ae72c8ff04e13ab3be7ae4861e1
3
+ size 4787616448
model-00004-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cdc3fde66f211aa2696a856bfe45fc64ed9e5b755a5b7b6ecf327167cae27f7b
3
+ size 4787616584
model-00005-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:59be449a6df2c10e8c4a606fc742f4280a965351241d11bf63b36b3cc1dc1d2d
3
+ size 4787616584
model-00006-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:9c65bffe8e7612964a601ed81d9031dad29436a2ab110446639d6ce84bda9478
3
+ size 4787616584
model-00007-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:955d38bf18ad91391302301af89a7b4b870c49223d327f1497a44d43bb1d8c21
3
+ size 4787616584
model-00008-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:954a7d1458d4f5d66777bd41ee6470dbda22cbce8f383a0c4650b72cf2a6b12a
3
+ size 4787616584
model-00009-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c244c4c6a3c5c25d3f562dc1df3423f034f60ac3e9992253539c8fce6d907de8
3
+ size 4787616584
model-00010-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d8d60bbf9a7292d3f32d73a060ce3f59220b5d5f198caa8e89c4736754636b18
3
+ size 4787616584
model-00011-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:39a9a46c00166ed544028e467f22cf16c86110152234efcdb994074a82e4a012
3
+ size 4787616584
model-00012-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c53065e01f6bdd73778908cfd23104dffaa1d80d9d25611719ff8fbc59fa2ca5
3
+ size 4787616584
model-00013-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e5bdb09db82110947c98d0542351d0c7f8353e9148e320e6ba60945824c388c1
3
+ size 4787616584
model-00014-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:81ccdcc1b6c554b508c4866b215493a0cf24d6ef280300f756eea705b5ffef69
3
+ size 4787616584
model-00015-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:df284aa2cccc6b849bd36579e4fa660d7f469ec9a998540ad11ac903dbd71ac6
3
+ size 4787616584
model-00016-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e46674e6a54ccecb95486d16dbcc1d1691227065be8c7108decba294c7cd3dd3
3
+ size 4787616584
model-00017-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0d6f5871cc5e09b3997a59abb33b284417c0434e407aa3ca2dc6d24a7c406fcb
3
+ size 4787616584
model-00018-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:baa2703cf4f297896c8ee1cdb2ca5f9ab6c2417d3a3f7e246cd9117efd7262c0
3
+ size 4787616712
model-00019-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:aa5403b9a1163578e6cd1293032c481568d66fdb174c0146e04e380caa94aba3
3
+ size 4787616848
model-00020-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3792e6b4ae571c7ff44357a0b5058c2fb636788ae614295e5409e166c99560f7
3
+ size 4787616848
model-00021-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:186cd6155bcbe4c2a746317d2a9d646b1c985cc9459a2120270e03cb2838bd17
3
+ size 4787616848
model-00022-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7c963f5a4fbdfc48b1845aa1dce68bf99d5df3e02a4722220a00b237f91767bd
3
+ size 4208055480
model-00023-of-00023.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:71b079a02069ebe5a5f4f8135815f136afb1d6fc4423620dcb2a14fa399cdf12
3
+ size 4202692736
model.safetensors.index.json ADDED
The diff for this file is too large to render. See raw diff