ashioyajotham commited on
Commit
4c15885
1 Parent(s): 261330e

falcon-coder

Browse files
README.md CHANGED
@@ -1,7 +1,10 @@
1
  ---
2
- base_model: ybelkada/falcon-7b-sharded-bf16
3
  tags:
 
 
4
  - generated_from_trainer
 
5
  model-index:
6
  - name: results
7
  results: []
@@ -49,7 +52,8 @@ The following hyperparameters were used during training:
49
 
50
  ### Framework versions
51
 
52
- - Transformers 4.35.2
53
- - Pytorch 2.1.0+cu118
54
- - Datasets 2.15.0
55
- - Tokenizers 0.15.0
 
 
1
  ---
2
+ library_name: peft
3
  tags:
4
+ - trl
5
+ - sft
6
  - generated_from_trainer
7
+ base_model: ybelkada/falcon-7b-sharded-bf16
8
  model-index:
9
  - name: results
10
  results: []
 
52
 
53
  ### Framework versions
54
 
55
+ - PEFT 0.8.2
56
+ - Transformers 4.37.2
57
+ - Pytorch 2.1.0+cu121
58
+ - Datasets 2.17.0
59
+ - Tokenizers 0.15.2
adapter_config.json CHANGED
@@ -19,10 +19,11 @@
19
  "rank_pattern": {},
20
  "revision": null,
21
  "target_modules": [
22
- "dense",
23
  "dense_h_to_4h",
24
- "query_key_value",
25
- "dense_4h_to_h"
26
  ],
27
- "task_type": "CAUSAL_LM"
 
28
  }
 
19
  "rank_pattern": {},
20
  "revision": null,
21
  "target_modules": [
22
+ "dense_4h_to_h",
23
  "dense_h_to_4h",
24
+ "dense",
25
+ "query_key_value"
26
  ],
27
+ "task_type": "CAUSAL_LM",
28
+ "use_rslora": false
29
  }
adapter_model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:f0e407e7782fb2a0bf5122dadf8d7d92014030e648107c4e3b7d4f91371b6258
3
  size 522227376
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:8a2ca6fda46f10f00ce92df30ef139db778d4d163f6c757202d0d25b8a2107eb
3
  size 522227376
runs/Feb19_09-35-08_92054d03d1d7/events.out.tfevents.1708335345.92054d03d1d7.310.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d6dbe9d06ed3e1a304a9432596ce0b0344f41609c90e99b1d3e063f8550e2101
3
+ size 6000
special_tokens_map.json CHANGED
@@ -12,6 +12,12 @@
12
  ">>SUFFIX<<",
13
  ">>MIDDLE<<"
14
  ],
15
- "eos_token": "<|endoftext|>",
 
 
 
 
 
 
16
  "pad_token": "<|endoftext|>"
17
  }
 
12
  ">>SUFFIX<<",
13
  ">>MIDDLE<<"
14
  ],
15
+ "eos_token": {
16
+ "content": "<|endoftext|>",
17
+ "lstrip": false,
18
+ "normalized": false,
19
+ "rstrip": false,
20
+ "single_word": false
21
+ },
22
  "pad_token": "<|endoftext|>"
23
  }
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:cb3a040691ecd3e262cc994a98409a144f5b10437b1c6c6c863b462342d1b623
3
- size 4600
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:66f07506ece4a41bb52b0b201ce1bddd16a2b319175b20499fe46b5a63fa8100
3
+ size 4728