[WIP] Upload folder using huggingface_hub (multi-commit 77778710e02cfd07373044489d4c6e5350ab070449633d3ef865e187ea581ac5)

#1
by luodian - opened
added_tokens.json ADDED
@@ -0,0 +1,5 @@
 
 
 
 
 
 
1
+ {
2
+ "<|endoftext|>": 151643,
3
+ "<|im_end|>": 151645,
4
+ "<|im_start|>": 151644
5
+ }
checkpoint-9000/added_tokens.json ADDED
@@ -0,0 +1,5 @@
 
 
 
 
 
 
1
+ {
2
+ "<|endoftext|>": 151643,
3
+ "<|im_end|>": 151645,
4
+ "<|im_start|>": 151644
5
+ }
checkpoint-9000/config.json ADDED
@@ -0,0 +1,197 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "/mnt/bn/vl-research/checkpoints/onevision/llavanext-google_siglip-so400m-patch14-384-Qwen_Qwen2-0.5B-Instruct-mid_to_final_next_3p2m_am9_july21",
3
+ "architectures": [
4
+ "LlavaQwenForCausalLM"
5
+ ],
6
+ "attention_dropout": 0.0,
7
+ "bos_token_id": 151643,
8
+ "eos_token_id": 151645,
9
+ "hidden_act": "silu",
10
+ "hidden_size": 896,
11
+ "image_aspect_ratio": "anyres_max_9",
12
+ "image_crop_resolution": null,
13
+ "image_grid_pinpoints": [
14
+ [
15
+ 384,
16
+ 384
17
+ ],
18
+ [
19
+ 384,
20
+ 768
21
+ ],
22
+ [
23
+ 384,
24
+ 1152
25
+ ],
26
+ [
27
+ 384,
28
+ 1536
29
+ ],
30
+ [
31
+ 384,
32
+ 1920
33
+ ],
34
+ [
35
+ 384,
36
+ 2304
37
+ ],
38
+ [
39
+ 768,
40
+ 384
41
+ ],
42
+ [
43
+ 768,
44
+ 768
45
+ ],
46
+ [
47
+ 768,
48
+ 1152
49
+ ],
50
+ [
51
+ 768,
52
+ 1536
53
+ ],
54
+ [
55
+ 768,
56
+ 1920
57
+ ],
58
+ [
59
+ 768,
60
+ 2304
61
+ ],
62
+ [
63
+ 1152,
64
+ 384
65
+ ],
66
+ [
67
+ 1152,
68
+ 768
69
+ ],
70
+ [
71
+ 1152,
72
+ 1152
73
+ ],
74
+ [
75
+ 1152,
76
+ 1536
77
+ ],
78
+ [
79
+ 1152,
80
+ 1920
81
+ ],
82
+ [
83
+ 1152,
84
+ 2304
85
+ ],
86
+ [
87
+ 1536,
88
+ 384
89
+ ],
90
+ [
91
+ 1536,
92
+ 768
93
+ ],
94
+ [
95
+ 1536,
96
+ 1152
97
+ ],
98
+ [
99
+ 1536,
100
+ 1536
101
+ ],
102
+ [
103
+ 1536,
104
+ 1920
105
+ ],
106
+ [
107
+ 1536,
108
+ 2304
109
+ ],
110
+ [
111
+ 1920,
112
+ 384
113
+ ],
114
+ [
115
+ 1920,
116
+ 768
117
+ ],
118
+ [
119
+ 1920,
120
+ 1152
121
+ ],
122
+ [
123
+ 1920,
124
+ 1536
125
+ ],
126
+ [
127
+ 1920,
128
+ 1920
129
+ ],
130
+ [
131
+ 1920,
132
+ 2304
133
+ ],
134
+ [
135
+ 2304,
136
+ 384
137
+ ],
138
+ [
139
+ 2304,
140
+ 768
141
+ ],
142
+ [
143
+ 2304,
144
+ 1152
145
+ ],
146
+ [
147
+ 2304,
148
+ 1536
149
+ ],
150
+ [
151
+ 2304,
152
+ 1920
153
+ ],
154
+ [
155
+ 2304,
156
+ 2304
157
+ ]
158
+ ],
159
+ "image_split_resolution": null,
160
+ "initializer_range": 0.02,
161
+ "intermediate_size": 4864,
162
+ "max_position_embeddings": 32768,
163
+ "max_window_layers": 24,
164
+ "mm_hidden_size": 1152,
165
+ "mm_patch_merge_type": "spatial_unpad",
166
+ "mm_projector_lr": null,
167
+ "mm_projector_type": "mlp2x_gelu",
168
+ "mm_resampler_type": null,
169
+ "mm_spatial_pool_mode": "bilinear",
170
+ "mm_tunable_parts": "mm_vision_tower,mm_mlp_adapter,mm_language_model",
171
+ "mm_use_im_patch_token": false,
172
+ "mm_use_im_start_end": false,
173
+ "mm_vision_select_feature": "patch",
174
+ "mm_vision_select_layer": -2,
175
+ "mm_vision_tower": "google/siglip-so400m-patch14-384",
176
+ "mm_vision_tower_lr": 2e-06,
177
+ "model_type": "qwen2",
178
+ "num_attention_heads": 14,
179
+ "num_hidden_layers": 24,
180
+ "num_key_value_heads": 2,
181
+ "pos_skipping_range": 4096,
182
+ "rms_norm_eps": 1e-06,
183
+ "rope_scaling": null,
184
+ "rope_theta": 1000000.0,
185
+ "sliding_window": 32768,
186
+ "tie_word_embeddings": true,
187
+ "tokenizer_model_max_length": 32768,
188
+ "tokenizer_padding_side": "right",
189
+ "torch_dtype": "bfloat16",
190
+ "transformers_version": "4.40.0.dev0",
191
+ "use_cache": false,
192
+ "use_mm_proj": true,
193
+ "use_pos_skipping": false,
194
+ "use_sliding_window": false,
195
+ "vision_tower_pretrained": null,
196
+ "vocab_size": 151936
197
+ }
checkpoint-9000/generation_config.json ADDED
@@ -0,0 +1,14 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token_id": 151643,
3
+ "do_sample": true,
4
+ "eos_token_id": [
5
+ 151645,
6
+ 151643
7
+ ],
8
+ "pad_token_id": 151643,
9
+ "repetition_penalty": 1.1,
10
+ "temperature": 0.7,
11
+ "top_k": 20,
12
+ "top_p": 0.8,
13
+ "transformers_version": "4.40.0.dev0"
14
+ }
checkpoint-9000/global_step9000/bf16_zero_pp_rank_0_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c1ebd3b3f804a4442119f73d2740ddd67099b809f6cdc35830e0b7020d152719
3
+ size 167561526
checkpoint-9000/global_step9000/bf16_zero_pp_rank_10_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2ce1939712beb6dce9fb71ceca700f9cea08d7a78792d45226f4a53b492e88f7
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_11_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3d8b4850999a3d9b9a4ea0c08a1695571fd544dda07a24bf87fe64dbf166d71c
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_12_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:bf8e7c21c49adcb43de1a14fbfad582361a4b807dc426e2d160997325c8fde3c
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_13_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e0c08be787c2fcb91e484cb57f98f84240bc94b47dc54c8bb22575638dec0c26
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_14_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:43da855ec2de285be0fbd86fc056ed2f665c62170a608dd26cdb688dfe7bbd17
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_15_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:bd2f56a43ea5199a7f056858ab6d1316aa2ec39baefefe6dc2fd72fadbb9d4ef
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_16_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:dec108d0fa4ab9f6f48b3d97d24d26cf29cb1b915c892906db02d6d016553625
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_17_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2b6a3271fd25ef5df044c9dc4c5a9826ca51e22893fd3a298eedaa03aaf0dfac
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_18_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:3d2a5b66df62749d05f6c6032c20e920955f9d159e0bb7cd0f275d6edfa6a213
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_19_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:90b1a040e30d2496cfef669efa67fc584b8c09bc8b16226d60835943a90ddfd1
3
+ size 167561546
checkpoint-9000/global_step9000/bf16_zero_pp_rank_1_mp_rank_00_optim_states.pt ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e40188464769400cf8a2f685ad909056d3a22bde54f7943c8d6447c27d0c5e82
3
+ size 167561526