taesiri commited on
Commit
57b9730
1 Parent(s): 2be73ff

Upload 15 files

Browse files
.gitattributes CHANGED
@@ -33,3 +33,5 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
 
 
 
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
36
+ train-set-60K.json filter=lfs diff=lfs merge=lfs -text
37
+ val-set-60K.json filter=lfs diff=lfs merge=lfs -text
README.md ADDED
@@ -0,0 +1,20 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ ---
2
+ library_name: peft
3
+ ---
4
+ ## Training procedure
5
+
6
+
7
+ The following `bitsandbytes` quantization config was used during training:
8
+ - load_in_8bit: False
9
+ - load_in_4bit: True
10
+ - llm_int8_threshold: 6.0
11
+ - llm_int8_skip_modules: ['mm_projector']
12
+ - llm_int8_enable_fp32_cpu_offload: False
13
+ - llm_int8_has_fp16_weight: False
14
+ - bnb_4bit_quant_type: nf4
15
+ - bnb_4bit_use_double_quant: True
16
+ - bnb_4bit_compute_dtype: bfloat16
17
+ ### Framework versions
18
+
19
+
20
+ - PEFT 0.4.0
adapter_config.json ADDED
@@ -0,0 +1,26 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "auto_mapping": null,
3
+ "base_model_name_or_path": "./temp/llava-v1.5-13b",
4
+ "bias": "none",
5
+ "fan_in_fan_out": false,
6
+ "inference_mode": true,
7
+ "init_lora_weights": true,
8
+ "layers_pattern": null,
9
+ "layers_to_transform": null,
10
+ "lora_alpha": 256,
11
+ "lora_dropout": 0.05,
12
+ "modules_to_save": null,
13
+ "peft_type": "LORA",
14
+ "r": 128,
15
+ "revision": null,
16
+ "target_modules": [
17
+ "k_proj",
18
+ "v_proj",
19
+ "down_proj",
20
+ "q_proj",
21
+ "gate_proj",
22
+ "up_proj",
23
+ "o_proj"
24
+ ],
25
+ "task_type": "CAUSAL_LM"
26
+ }
adapter_model.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:51fcae5e69e19af5efb2e5860bc8ffaa756162359bb1674a803d07044bc3f59d
3
+ size 1001584333
best_llava_eval_model/README.md ADDED
@@ -0,0 +1,56 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ ---
2
+ library_name: peft
3
+ ---
4
+ ## Training procedure
5
+
6
+
7
+ The following `bitsandbytes` quantization config was used during training:
8
+ - load_in_8bit: False
9
+ - load_in_4bit: True
10
+ - llm_int8_threshold: 6.0
11
+ - llm_int8_skip_modules: ['mm_projector']
12
+ - llm_int8_enable_fp32_cpu_offload: False
13
+ - llm_int8_has_fp16_weight: False
14
+ - bnb_4bit_quant_type: nf4
15
+ - bnb_4bit_use_double_quant: True
16
+ - bnb_4bit_compute_dtype: bfloat16
17
+
18
+ The following `bitsandbytes` quantization config was used during training:
19
+ - load_in_8bit: False
20
+ - load_in_4bit: True
21
+ - llm_int8_threshold: 6.0
22
+ - llm_int8_skip_modules: ['mm_projector']
23
+ - llm_int8_enable_fp32_cpu_offload: False
24
+ - llm_int8_has_fp16_weight: False
25
+ - bnb_4bit_quant_type: nf4
26
+ - bnb_4bit_use_double_quant: True
27
+ - bnb_4bit_compute_dtype: bfloat16
28
+
29
+ The following `bitsandbytes` quantization config was used during training:
30
+ - load_in_8bit: False
31
+ - load_in_4bit: True
32
+ - llm_int8_threshold: 6.0
33
+ - llm_int8_skip_modules: ['mm_projector']
34
+ - llm_int8_enable_fp32_cpu_offload: False
35
+ - llm_int8_has_fp16_weight: False
36
+ - bnb_4bit_quant_type: nf4
37
+ - bnb_4bit_use_double_quant: True
38
+ - bnb_4bit_compute_dtype: bfloat16
39
+
40
+ The following `bitsandbytes` quantization config was used during training:
41
+ - load_in_8bit: False
42
+ - load_in_4bit: True
43
+ - llm_int8_threshold: 6.0
44
+ - llm_int8_skip_modules: ['mm_projector']
45
+ - llm_int8_enable_fp32_cpu_offload: False
46
+ - llm_int8_has_fp16_weight: False
47
+ - bnb_4bit_quant_type: nf4
48
+ - bnb_4bit_use_double_quant: True
49
+ - bnb_4bit_compute_dtype: bfloat16
50
+ ### Framework versions
51
+
52
+ - PEFT 0.4.0
53
+ - PEFT 0.4.0
54
+ - PEFT 0.4.0
55
+
56
+ - PEFT 0.4.0
best_llava_eval_model/adapter_config.json ADDED
@@ -0,0 +1,26 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "auto_mapping": null,
3
+ "base_model_name_or_path": "./temp/llava-v1.5-13b",
4
+ "bias": "none",
5
+ "fan_in_fan_out": false,
6
+ "inference_mode": true,
7
+ "init_lora_weights": true,
8
+ "layers_pattern": null,
9
+ "layers_to_transform": null,
10
+ "lora_alpha": 256,
11
+ "lora_dropout": 0.05,
12
+ "modules_to_save": null,
13
+ "peft_type": "LORA",
14
+ "r": 128,
15
+ "revision": null,
16
+ "target_modules": [
17
+ "k_proj",
18
+ "v_proj",
19
+ "down_proj",
20
+ "q_proj",
21
+ "gate_proj",
22
+ "up_proj",
23
+ "o_proj"
24
+ ],
25
+ "task_type": "CAUSAL_LM"
26
+ }
best_llava_eval_model/adapter_model.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d5c62599526e77831d1349b45872698d2d0a8f7765dbab4029319c231a7e06fd
3
+ size 1001584333
best_llava_eval_model/config.json ADDED
@@ -0,0 +1,58 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "./temp/llava-v1.5-13b",
3
+ "architectures": [
4
+ "LlavaLlamaForCausalLM"
5
+ ],
6
+ "bos_token_id": 1,
7
+ "eos_token_id": 2,
8
+ "freeze_mm_mlp_adapter": false,
9
+ "freeze_mm_vision_resampler": false,
10
+ "hidden_act": "silu",
11
+ "hidden_size": 5120,
12
+ "image_aspect_ratio": "pad",
13
+ "initializer_range": 0.02,
14
+ "intermediate_size": 13824,
15
+ "max_length": 4096,
16
+ "max_position_embeddings": 4096,
17
+ "mm_hidden_size": 1024,
18
+ "mm_projector_lr": 2e-05,
19
+ "mm_projector_type": "mlp2x_gelu",
20
+ "mm_resampler_type": null,
21
+ "mm_use_im_patch_token": false,
22
+ "mm_use_im_start_end": false,
23
+ "mm_vision_select_feature": "patch",
24
+ "mm_vision_select_layer": -2,
25
+ "mm_vision_tower": "openai/clip-vit-large-patch14-336",
26
+ "model_type": "llava",
27
+ "num_attention_heads": 40,
28
+ "num_hidden_layers": 40,
29
+ "num_key_value_heads": 40,
30
+ "pad_token_id": 0,
31
+ "pretraining_tp": 1,
32
+ "quantization_config": {
33
+ "bnb_4bit_compute_dtype": "bfloat16",
34
+ "bnb_4bit_quant_type": "nf4",
35
+ "bnb_4bit_use_double_quant": true,
36
+ "llm_int8_enable_fp32_cpu_offload": false,
37
+ "llm_int8_has_fp16_weight": false,
38
+ "llm_int8_skip_modules": [
39
+ "mm_projector"
40
+ ],
41
+ "llm_int8_threshold": 6.0,
42
+ "load_in_4bit": true,
43
+ "load_in_8bit": false
44
+ },
45
+ "rms_norm_eps": 1e-05,
46
+ "rope_scaling": null,
47
+ "tie_word_embeddings": false,
48
+ "tokenizer_model_max_length": 2048,
49
+ "tokenizer_padding_side": "right",
50
+ "torch_dtype": "bfloat16",
51
+ "transformers_version": "4.31.0",
52
+ "tune_mm_mlp_adapter": false,
53
+ "tune_mm_vision_resampler": false,
54
+ "unfreeze_mm_vision_tower": false,
55
+ "use_cache": false,
56
+ "use_mm_proj": true,
57
+ "vocab_size": 32000
58
+ }
best_llava_eval_model/mm_projector.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cef35115bca05e7ab8d4c3ad575b2acfca5999beb2801145f8f518beaa606d3b
3
+ size 62936701
best_llava_eval_model/non_lora_trainables.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d3565188272471b3d56f17a750aedc815333bcff9cb7e99a9953a22b7af946d4
3
+ size 62936807
config.json ADDED
@@ -0,0 +1,58 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "./temp/llava-v1.5-13b",
3
+ "architectures": [
4
+ "LlavaLlamaForCausalLM"
5
+ ],
6
+ "bos_token_id": 1,
7
+ "eos_token_id": 2,
8
+ "freeze_mm_mlp_adapter": false,
9
+ "freeze_mm_vision_resampler": false,
10
+ "hidden_act": "silu",
11
+ "hidden_size": 5120,
12
+ "image_aspect_ratio": "pad",
13
+ "initializer_range": 0.02,
14
+ "intermediate_size": 13824,
15
+ "max_length": 4096,
16
+ "max_position_embeddings": 4096,
17
+ "mm_hidden_size": 1024,
18
+ "mm_projector_lr": 2e-05,
19
+ "mm_projector_type": "mlp2x_gelu",
20
+ "mm_resampler_type": null,
21
+ "mm_use_im_patch_token": false,
22
+ "mm_use_im_start_end": false,
23
+ "mm_vision_select_feature": "patch",
24
+ "mm_vision_select_layer": -2,
25
+ "mm_vision_tower": "openai/clip-vit-large-patch14-336",
26
+ "model_type": "llava",
27
+ "num_attention_heads": 40,
28
+ "num_hidden_layers": 40,
29
+ "num_key_value_heads": 40,
30
+ "pad_token_id": 0,
31
+ "pretraining_tp": 1,
32
+ "quantization_config": {
33
+ "bnb_4bit_compute_dtype": "bfloat16",
34
+ "bnb_4bit_quant_type": "nf4",
35
+ "bnb_4bit_use_double_quant": true,
36
+ "llm_int8_enable_fp32_cpu_offload": false,
37
+ "llm_int8_has_fp16_weight": false,
38
+ "llm_int8_skip_modules": [
39
+ "mm_projector"
40
+ ],
41
+ "llm_int8_threshold": 6.0,
42
+ "load_in_4bit": true,
43
+ "load_in_8bit": false
44
+ },
45
+ "rms_norm_eps": 1e-05,
46
+ "rope_scaling": null,
47
+ "tie_word_embeddings": false,
48
+ "tokenizer_model_max_length": 2048,
49
+ "tokenizer_padding_side": "right",
50
+ "torch_dtype": "bfloat16",
51
+ "transformers_version": "4.31.0",
52
+ "tune_mm_mlp_adapter": false,
53
+ "tune_mm_vision_resampler": false,
54
+ "unfreeze_mm_vision_tower": false,
55
+ "use_cache": true,
56
+ "use_mm_proj": true,
57
+ "vocab_size": 32000
58
+ }
mm_projector.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cef35115bca05e7ab8d4c3ad575b2acfca5999beb2801145f8f518beaa606d3b
3
+ size 62936701
non_lora_trainables.bin ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:19ed1efd1b56a90fc87144b930e4ba8f96f4c37fdca60b8ce25c903f5795ed5e
3
+ size 62936807
train-set-60K.json ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a325c01b903f05f017fe218401cefa32a2bf0373cd60a6515cdbdfcf457420bc
3
+ size 192426529
trainer_state.json ADDED
The diff for this file is too large to render. See raw diff
 
val-set-60K.json ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:23fcbd741624613268a06cb5dcfe3b1978e55f1b00c904829910bae5f2c43e8e
3
+ size 48003730