Training in progress, step 100
Browse files- arcade100k.tiktoken +0 -0
- config.json +33 -0
- model.safetensors +3 -0
- runs/Feb16_23-12-10_cccxc540/events.out.tfevents.1708143188.cccxc540.4000843.0 +3 -0
- runs/Feb16_23-16-03_cccxc540/events.out.tfevents.1708143390.cccxc540.4001861.0 +3 -0
- runs/Feb16_23-18-11_cccxc540/events.out.tfevents.1708143518.cccxc540.4002534.0 +3 -0
- runs/Feb16_23-26-28_cccxc540/events.out.tfevents.1708144018.cccxc540.4004223.0 +3 -0
- runs/Feb16_23-28-09_cccxc540/events.out.tfevents.1708144118.cccxc540.4004962.0 +3 -0
- runs/Feb16_23-51-16_cccxc540/events.out.tfevents.1708145504.cccxc540.4008629.0 +3 -0
- runs/Feb19_22-42-48_cccxc544/events.out.tfevents.1708400646.cccxc544.374957.0 +3 -0
- runs/Feb19_22-54-24_cccxc544/events.out.tfevents.1708401285.cccxc544.377003.0 +3 -0
- runs/Feb19_22-57-38_cccxc544/events.out.tfevents.1708401478.cccxc544.377500.0 +3 -0
- runs/Feb19_22-58-44_cccxc544/events.out.tfevents.1708401544.cccxc544.377685.0 +3 -0
- runs/Feb19_22-59-31_cccxc544/events.out.tfevents.1708401592.cccxc544.377856.0 +3 -0
- runs/Feb19_23-00-27_cccxc544/events.out.tfevents.1708401645.cccxc544.378225.0 +3 -0
- runs/Feb19_23-02-25_cccxc544/events.out.tfevents.1708401776.cccxc544.378452.0 +3 -0
- runs/Feb19_23-08-06_cccxc542/events.out.tfevents.1708402191.cccxc542.93482.0 +3 -0
- special_tokens_map.json +5 -0
- tokenizer_config.json +17 -0
- training_args.bin +3 -0
arcade100k.tiktoken
ADDED
The diff for this file is too large to render.
See raw diff
|
|
config.json
ADDED
@@ -0,0 +1,33 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"_name_or_path": "stabilityai/stablelm-2-zephyr-1_6b",
|
3 |
+
"architectures": [
|
4 |
+
"StableLMEpochForCausalLM"
|
5 |
+
],
|
6 |
+
"attention_dropout": 0.0,
|
7 |
+
"auto_map": {
|
8 |
+
"AutoConfig": "configuration_stablelm_epoch.StableLMEpochConfig",
|
9 |
+
"AutoModelForCausalLM": "modeling_stablelm_epoch.StableLMEpochForCausalLM"
|
10 |
+
},
|
11 |
+
"bos_token_id": 100257,
|
12 |
+
"eos_token_id": 100257,
|
13 |
+
"hidden_act": "silu",
|
14 |
+
"hidden_size": 2048,
|
15 |
+
"initializer_range": 0.02,
|
16 |
+
"intermediate_size": 5632,
|
17 |
+
"max_position_embeddings": 4096,
|
18 |
+
"model_type": "stablelm_epoch",
|
19 |
+
"norm_eps": 1e-05,
|
20 |
+
"num_attention_heads": 32,
|
21 |
+
"num_heads": 32,
|
22 |
+
"num_hidden_layers": 24,
|
23 |
+
"num_key_value_heads": 32,
|
24 |
+
"rope_pct": 0.25,
|
25 |
+
"rope_theta": 10000,
|
26 |
+
"rotary_scaling_factor": 1.0,
|
27 |
+
"tie_word_embeddings": false,
|
28 |
+
"torch_dtype": "bfloat16",
|
29 |
+
"transformers_version": "4.36.2",
|
30 |
+
"use_cache": false,
|
31 |
+
"use_qkv_bias": true,
|
32 |
+
"vocab_size": 100352
|
33 |
+
}
|
model.safetensors
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:68fd4b889e4c185bf5e5e3a59a64ec59f4715b90205c5ae1098c2c6964442dfd
|
3 |
+
size 3289069520
|
runs/Feb16_23-12-10_cccxc540/events.out.tfevents.1708143188.cccxc540.4000843.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:758161e2b7ebd2ab9dd6be5521d31133043c7f3183ce10cd60ca739fce29266f
|
3 |
+
size 4815
|
runs/Feb16_23-16-03_cccxc540/events.out.tfevents.1708143390.cccxc540.4001861.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b4402f643fbcf09c9dbb671a7d4e26993e4c3d729d4e9792217c2d6035bebd2a
|
3 |
+
size 6061
|
runs/Feb16_23-18-11_cccxc540/events.out.tfevents.1708143518.cccxc540.4002534.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a00466eade3ff3f8829b8982c81864c3abdfd6a696b829e9018fd8b590242dc8
|
3 |
+
size 4815
|
runs/Feb16_23-26-28_cccxc540/events.out.tfevents.1708144018.cccxc540.4004223.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:d7b7835dfdce0d4bff240a93b00de790bb013fd1f493204381aa1f51af6d7c4e
|
3 |
+
size 5438
|
runs/Feb16_23-28-09_cccxc540/events.out.tfevents.1708144118.cccxc540.4004962.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:105496165fcba8aea97f41e798fbe670fa669d22402acdb2b281717fb864c28f
|
3 |
+
size 5438
|
runs/Feb16_23-51-16_cccxc540/events.out.tfevents.1708145504.cccxc540.4008629.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ac65ff6e77861825bbe56986a698cdccde583c70993bea71bd1ee73eba5c5b67
|
3 |
+
size 9176
|
runs/Feb19_22-42-48_cccxc544/events.out.tfevents.1708400646.cccxc544.374957.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:017fc844be050cfc8e8ac069a3189ab91e0168bc86ed510321955aad90601605
|
3 |
+
size 5438
|
runs/Feb19_22-54-24_cccxc544/events.out.tfevents.1708401285.cccxc544.377003.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5819e0a1f81dc4bef9cd081af082d43beb225e82d98f025dca6aeb85b14ef9af
|
3 |
+
size 4817
|
runs/Feb19_22-57-38_cccxc544/events.out.tfevents.1708401478.cccxc544.377500.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1a6bf2f1c5a45557181b9ad9cd2886f261430915d28cef9300b334f2155a1a29
|
3 |
+
size 4817
|
runs/Feb19_22-58-44_cccxc544/events.out.tfevents.1708401544.cccxc544.377685.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:6d0afab6104d1781fa15c90e46dfddae1e386e27fa2cc0cacdd4ff105c8fc622
|
3 |
+
size 4817
|
runs/Feb19_22-59-31_cccxc544/events.out.tfevents.1708401592.cccxc544.377856.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:3f63edaffc6fba62f70edea4269f771557d9477693bb66a651f22a25292ce268
|
3 |
+
size 4817
|
runs/Feb19_23-00-27_cccxc544/events.out.tfevents.1708401645.cccxc544.378225.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:aea0a237e55de4769d65e1b179a29ed773df54a8af83fe7b93d6c4ba0c6ba2da
|
3 |
+
size 4817
|
runs/Feb19_23-02-25_cccxc544/events.out.tfevents.1708401776.cccxc544.378452.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b945cb0c22ca32173f7e1846b8b625b6592ba4bc4ddf7a23a237eb3aeb37380c
|
3 |
+
size 6686
|
runs/Feb19_23-08-06_cccxc542/events.out.tfevents.1708402191.cccxc542.93482.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:dfd76add1116e97fa777ff8147db8730c32919b4d2d320fe0656b92adef80eb8
|
3 |
+
size 12397
|
special_tokens_map.json
ADDED
@@ -0,0 +1,5 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"bos_token": "<|endoftext|>",
|
3 |
+
"eos_token": "<|endoftext|>",
|
4 |
+
"pad_token": "<|endoftext|>"
|
5 |
+
}
|
tokenizer_config.json
ADDED
@@ -0,0 +1,17 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"added_tokens_decoder": {},
|
3 |
+
"auto_map": {
|
4 |
+
"AutoTokenizer": [
|
5 |
+
"stabilityai/stablelm-2-zephyr-1_6b--tokenization_arcade100k.Arcade100kTokenizer",
|
6 |
+
null
|
7 |
+
]
|
8 |
+
},
|
9 |
+
"bos_token": "<|endoftext|>",
|
10 |
+
"chat_template": "{% for message in messages %}\n{% if message['role'] == 'user' %}\n{{ '<|user|>\n' + message['content'] + eos_token }}\n{% elif message['role'] == 'system' %}\n{{ '<|system|>\n' + message['content'] + eos_token }}\n{% elif message['role'] == 'assistant' %}\n{{ '<|assistant|>\n' + message['content'] + eos_token }}\n{% endif %}\n{% if loop.last and add_generation_prompt %}\n{{ '<|assistant|>' }}\n{% endif %}\n{% endfor %}",
|
11 |
+
"clean_up_tokenization_spaces": true,
|
12 |
+
"eos_token": "<|endoftext|>",
|
13 |
+
"errors": "replace",
|
14 |
+
"model_max_length": 2048,
|
15 |
+
"pad_token": "<|endoftext|>",
|
16 |
+
"tokenizer_class": "Arcade100kTokenizer"
|
17 |
+
}
|
training_args.bin
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:2791b5f6ccebcb20559d3e2e647d03fb6056e20c8301d57f11e4e205d97e9a26
|
3 |
+
size 5944
|