imelnyk commited on
Commit
a743e19
1 Parent(s): 73b1014

Training in progress, step 100

Browse files
arcade100k.tiktoken ADDED
The diff for this file is too large to render. See raw diff
 
config.json ADDED
@@ -0,0 +1,33 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "stabilityai/stablelm-2-zephyr-1_6b",
3
+ "architectures": [
4
+ "StableLMEpochForCausalLM"
5
+ ],
6
+ "attention_dropout": 0.0,
7
+ "auto_map": {
8
+ "AutoConfig": "configuration_stablelm_epoch.StableLMEpochConfig",
9
+ "AutoModelForCausalLM": "modeling_stablelm_epoch.StableLMEpochForCausalLM"
10
+ },
11
+ "bos_token_id": 100257,
12
+ "eos_token_id": 100257,
13
+ "hidden_act": "silu",
14
+ "hidden_size": 2048,
15
+ "initializer_range": 0.02,
16
+ "intermediate_size": 5632,
17
+ "max_position_embeddings": 4096,
18
+ "model_type": "stablelm_epoch",
19
+ "norm_eps": 1e-05,
20
+ "num_attention_heads": 32,
21
+ "num_heads": 32,
22
+ "num_hidden_layers": 24,
23
+ "num_key_value_heads": 32,
24
+ "rope_pct": 0.25,
25
+ "rope_theta": 10000,
26
+ "rotary_scaling_factor": 1.0,
27
+ "tie_word_embeddings": false,
28
+ "torch_dtype": "bfloat16",
29
+ "transformers_version": "4.36.2",
30
+ "use_cache": false,
31
+ "use_qkv_bias": true,
32
+ "vocab_size": 100352
33
+ }
model.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:68fd4b889e4c185bf5e5e3a59a64ec59f4715b90205c5ae1098c2c6964442dfd
3
+ size 3289069520
runs/Feb16_23-12-10_cccxc540/events.out.tfevents.1708143188.cccxc540.4000843.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:758161e2b7ebd2ab9dd6be5521d31133043c7f3183ce10cd60ca739fce29266f
3
+ size 4815
runs/Feb16_23-16-03_cccxc540/events.out.tfevents.1708143390.cccxc540.4001861.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:b4402f643fbcf09c9dbb671a7d4e26993e4c3d729d4e9792217c2d6035bebd2a
3
+ size 6061
runs/Feb16_23-18-11_cccxc540/events.out.tfevents.1708143518.cccxc540.4002534.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a00466eade3ff3f8829b8982c81864c3abdfd6a696b829e9018fd8b590242dc8
3
+ size 4815
runs/Feb16_23-26-28_cccxc540/events.out.tfevents.1708144018.cccxc540.4004223.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d7b7835dfdce0d4bff240a93b00de790bb013fd1f493204381aa1f51af6d7c4e
3
+ size 5438
runs/Feb16_23-28-09_cccxc540/events.out.tfevents.1708144118.cccxc540.4004962.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:105496165fcba8aea97f41e798fbe670fa669d22402acdb2b281717fb864c28f
3
+ size 5438
runs/Feb16_23-51-16_cccxc540/events.out.tfevents.1708145504.cccxc540.4008629.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ac65ff6e77861825bbe56986a698cdccde583c70993bea71bd1ee73eba5c5b67
3
+ size 9176
runs/Feb19_22-42-48_cccxc544/events.out.tfevents.1708400646.cccxc544.374957.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:017fc844be050cfc8e8ac069a3189ab91e0168bc86ed510321955aad90601605
3
+ size 5438
runs/Feb19_22-54-24_cccxc544/events.out.tfevents.1708401285.cccxc544.377003.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5819e0a1f81dc4bef9cd081af082d43beb225e82d98f025dca6aeb85b14ef9af
3
+ size 4817
runs/Feb19_22-57-38_cccxc544/events.out.tfevents.1708401478.cccxc544.377500.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:1a6bf2f1c5a45557181b9ad9cd2886f261430915d28cef9300b334f2155a1a29
3
+ size 4817
runs/Feb19_22-58-44_cccxc544/events.out.tfevents.1708401544.cccxc544.377685.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:6d0afab6104d1781fa15c90e46dfddae1e386e27fa2cc0cacdd4ff105c8fc622
3
+ size 4817
runs/Feb19_22-59-31_cccxc544/events.out.tfevents.1708401592.cccxc544.377856.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3f63edaffc6fba62f70edea4269f771557d9477693bb66a651f22a25292ce268
3
+ size 4817
runs/Feb19_23-00-27_cccxc544/events.out.tfevents.1708401645.cccxc544.378225.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:aea0a237e55de4769d65e1b179a29ed773df54a8af83fe7b93d6c4ba0c6ba2da
3
+ size 4817
runs/Feb19_23-02-25_cccxc544/events.out.tfevents.1708401776.cccxc544.378452.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:b945cb0c22ca32173f7e1846b8b625b6592ba4bc4ddf7a23a237eb3aeb37380c
3
+ size 6686
runs/Feb19_23-08-06_cccxc542/events.out.tfevents.1708402191.cccxc542.93482.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:dfd76add1116e97fa777ff8147db8730c32919b4d2d320fe0656b92adef80eb8
3
+ size 12397
special_tokens_map.json ADDED
@@ -0,0 +1,5 @@
 
 
 
 
 
 
1
+ {
2
+ "bos_token": "<|endoftext|>",
3
+ "eos_token": "<|endoftext|>",
4
+ "pad_token": "<|endoftext|>"
5
+ }
tokenizer_config.json ADDED
@@ -0,0 +1,17 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "added_tokens_decoder": {},
3
+ "auto_map": {
4
+ "AutoTokenizer": [
5
+ "stabilityai/stablelm-2-zephyr-1_6b--tokenization_arcade100k.Arcade100kTokenizer",
6
+ null
7
+ ]
8
+ },
9
+ "bos_token": "<|endoftext|>",
10
+ "chat_template": "{% for message in messages %}\n{% if message['role'] == 'user' %}\n{{ '<|user|>\n' + message['content'] + eos_token }}\n{% elif message['role'] == 'system' %}\n{{ '<|system|>\n' + message['content'] + eos_token }}\n{% elif message['role'] == 'assistant' %}\n{{ '<|assistant|>\n' + message['content'] + eos_token }}\n{% endif %}\n{% if loop.last and add_generation_prompt %}\n{{ '<|assistant|>' }}\n{% endif %}\n{% endfor %}",
11
+ "clean_up_tokenization_spaces": true,
12
+ "eos_token": "<|endoftext|>",
13
+ "errors": "replace",
14
+ "model_max_length": 2048,
15
+ "pad_token": "<|endoftext|>",
16
+ "tokenizer_class": "Arcade100kTokenizer"
17
+ }
training_args.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2791b5f6ccebcb20559d3e2e647d03fb6056e20c8301d57f11e4e205d97e9a26
3
+ size 5944