Upload processor
Browse files- added_tokens.json +75 -77
- tokenizer.json +74 -101
- tokenizer_config.json +1 -1
added_tokens.json
CHANGED
@@ -1,82 +1,80 @@
|
|
1 |
{
|
2 |
-
"''": 57526,
|
3 |
-
"1": 57525,
|
4 |
"<s_iitcdip>": 57523,
|
5 |
"<s_synthdog>": 57524,
|
6 |
"<sep/>": 57522,
|
7 |
-
"a''1 ":
|
8 |
-
"a''16 ":
|
9 |
-
"a''2 ":
|
10 |
-
"a''4 ":
|
11 |
-
"a''8 ":
|
12 |
-
"a'1 ":
|
13 |
-
"a'16 ":
|
14 |
-
"a'2 ":
|
15 |
-
"a'4 ":
|
16 |
-
"a'8 ":
|
17 |
-
"b''1 ":
|
18 |
-
"b''16 ":
|
19 |
-
"b''2 ":
|
20 |
-
"b''4 ":
|
21 |
-
"b''8 ":
|
22 |
-
"b'1 ":
|
23 |
-
"b'16 ":
|
24 |
-
"b'2 ":
|
25 |
-
"b'4 ":
|
26 |
-
"b'8 ":
|
27 |
-
"c''1 ":
|
28 |
-
"c''16 ":
|
29 |
-
"c''2 ":
|
30 |
-
"c''4 ":
|
31 |
-
"c''8 ":
|
32 |
-
"c'1 ":
|
33 |
-
"c'16 ":
|
34 |
-
"c'2 ":
|
35 |
-
"c'4 ":
|
36 |
-
"c'8 ":
|
37 |
-
"d''1 ":
|
38 |
-
"d''16 ":
|
39 |
-
"d''2 ":
|
40 |
-
"d''4 ":
|
41 |
-
"d''8 ":
|
42 |
-
"d'1 ":
|
43 |
-
"d'16 ":
|
44 |
-
"d'2 ":
|
45 |
-
"d'4 ":
|
46 |
-
"d'8 ":
|
47 |
-
"e''1 ":
|
48 |
-
"e''16 ":
|
49 |
-
"e''2 ":
|
50 |
-
"e''4 ":
|
51 |
-
"e''8 ":
|
52 |
-
"e'1 ":
|
53 |
-
"e'16 ":
|
54 |
-
"e'2 ":
|
55 |
-
"e'4 ":
|
56 |
-
"e'8 ":
|
57 |
-
"f''1 ":
|
58 |
-
"f''16 ":
|
59 |
-
"f''2 ":
|
60 |
-
"f''4 ":
|
61 |
-
"f''8 ":
|
62 |
-
"f'1 ":
|
63 |
-
"f'16 ":
|
64 |
-
"f'2 ":
|
65 |
-
"f'4 ":
|
66 |
-
"f'8 ":
|
67 |
-
"g''1 ":
|
68 |
-
"g''16 ":
|
69 |
-
"g''2 ":
|
70 |
-
"g''4 ":
|
71 |
-
"g''8 ":
|
72 |
-
"g'1 ":
|
73 |
-
"g'16 ":
|
74 |
-
"g'2 ":
|
75 |
-
"g'4 ":
|
76 |
-
"g'8 ":
|
77 |
-
"r1 ":
|
78 |
-
"r16 ":
|
79 |
-
"r2 ":
|
80 |
-
"r4 ":
|
81 |
-
"r8 ":
|
82 |
}
|
|
|
1 |
{
|
|
|
|
|
2 |
"<s_iitcdip>": 57523,
|
3 |
"<s_synthdog>": 57524,
|
4 |
"<sep/>": 57522,
|
5 |
+
"a''1 ": 57580,
|
6 |
+
"a''16 ": 57584,
|
7 |
+
"a''2 ": 57581,
|
8 |
+
"a''4 ": 57582,
|
9 |
+
"a''8 ": 57583,
|
10 |
+
"a'1 ": 57575,
|
11 |
+
"a'16 ": 57579,
|
12 |
+
"a'2 ": 57576,
|
13 |
+
"a'4 ": 57577,
|
14 |
+
"a'8 ": 57578,
|
15 |
+
"b''1 ": 57590,
|
16 |
+
"b''16 ": 57594,
|
17 |
+
"b''2 ": 57591,
|
18 |
+
"b''4 ": 57592,
|
19 |
+
"b''8 ": 57593,
|
20 |
+
"b'1 ": 57585,
|
21 |
+
"b'16 ": 57589,
|
22 |
+
"b'2 ": 57586,
|
23 |
+
"b'4 ": 57587,
|
24 |
+
"b'8 ": 57588,
|
25 |
+
"c''1 ": 57530,
|
26 |
+
"c''16 ": 57534,
|
27 |
+
"c''2 ": 57531,
|
28 |
+
"c''4 ": 57532,
|
29 |
+
"c''8 ": 57533,
|
30 |
+
"c'1 ": 57525,
|
31 |
+
"c'16 ": 57529,
|
32 |
+
"c'2 ": 57526,
|
33 |
+
"c'4 ": 57527,
|
34 |
+
"c'8 ": 57528,
|
35 |
+
"d''1 ": 57540,
|
36 |
+
"d''16 ": 57544,
|
37 |
+
"d''2 ": 57541,
|
38 |
+
"d''4 ": 57542,
|
39 |
+
"d''8 ": 57543,
|
40 |
+
"d'1 ": 57535,
|
41 |
+
"d'16 ": 57539,
|
42 |
+
"d'2 ": 57536,
|
43 |
+
"d'4 ": 57537,
|
44 |
+
"d'8 ": 57538,
|
45 |
+
"e''1 ": 57550,
|
46 |
+
"e''16 ": 57554,
|
47 |
+
"e''2 ": 57551,
|
48 |
+
"e''4 ": 57552,
|
49 |
+
"e''8 ": 57553,
|
50 |
+
"e'1 ": 57545,
|
51 |
+
"e'16 ": 57549,
|
52 |
+
"e'2 ": 57546,
|
53 |
+
"e'4 ": 57547,
|
54 |
+
"e'8 ": 57548,
|
55 |
+
"f''1 ": 57560,
|
56 |
+
"f''16 ": 57564,
|
57 |
+
"f''2 ": 57561,
|
58 |
+
"f''4 ": 57562,
|
59 |
+
"f''8 ": 57563,
|
60 |
+
"f'1 ": 57555,
|
61 |
+
"f'16 ": 57559,
|
62 |
+
"f'2 ": 57556,
|
63 |
+
"f'4 ": 57557,
|
64 |
+
"f'8 ": 57558,
|
65 |
+
"g''1 ": 57570,
|
66 |
+
"g''16 ": 57574,
|
67 |
+
"g''2 ": 57571,
|
68 |
+
"g''4 ": 57572,
|
69 |
+
"g''8 ": 57573,
|
70 |
+
"g'1 ": 57565,
|
71 |
+
"g'16 ": 57569,
|
72 |
+
"g'2 ": 57566,
|
73 |
+
"g'4 ": 57567,
|
74 |
+
"g'8 ": 57568,
|
75 |
+
"r1 ": 57595,
|
76 |
+
"r16 ": 57599,
|
77 |
+
"r2 ": 57596,
|
78 |
+
"r4 ": 57597,
|
79 |
+
"r8 ": 57598
|
80 |
}
|
tokenizer.json
CHANGED
@@ -39,15 +39,6 @@
|
|
39 |
"normalized": false,
|
40 |
"special": true
|
41 |
},
|
42 |
-
{
|
43 |
-
"id": 28431,
|
44 |
-
"content": "'",
|
45 |
-
"single_word": false,
|
46 |
-
"lstrip": false,
|
47 |
-
"rstrip": false,
|
48 |
-
"normalized": true,
|
49 |
-
"special": false
|
50 |
-
},
|
51 |
{
|
52 |
"id": 57521,
|
53 |
"content": "<mask>",
|
@@ -86,24 +77,6 @@
|
|
86 |
},
|
87 |
{
|
88 |
"id": 57525,
|
89 |
-
"content": "1",
|
90 |
-
"single_word": false,
|
91 |
-
"lstrip": false,
|
92 |
-
"rstrip": false,
|
93 |
-
"normalized": true,
|
94 |
-
"special": false
|
95 |
-
},
|
96 |
-
{
|
97 |
-
"id": 57526,
|
98 |
-
"content": "''",
|
99 |
-
"single_word": false,
|
100 |
-
"lstrip": false,
|
101 |
-
"rstrip": false,
|
102 |
-
"normalized": true,
|
103 |
-
"special": false
|
104 |
-
},
|
105 |
-
{
|
106 |
-
"id": 57527,
|
107 |
"content": "c'1 ",
|
108 |
"single_word": false,
|
109 |
"lstrip": false,
|
@@ -112,7 +85,7 @@
|
|
112 |
"special": false
|
113 |
},
|
114 |
{
|
115 |
-
"id":
|
116 |
"content": "c'2 ",
|
117 |
"single_word": false,
|
118 |
"lstrip": false,
|
@@ -121,7 +94,7 @@
|
|
121 |
"special": false
|
122 |
},
|
123 |
{
|
124 |
-
"id":
|
125 |
"content": "c'4 ",
|
126 |
"single_word": false,
|
127 |
"lstrip": false,
|
@@ -130,7 +103,7 @@
|
|
130 |
"special": false
|
131 |
},
|
132 |
{
|
133 |
-
"id":
|
134 |
"content": "c'8 ",
|
135 |
"single_word": false,
|
136 |
"lstrip": false,
|
@@ -139,7 +112,7 @@
|
|
139 |
"special": false
|
140 |
},
|
141 |
{
|
142 |
-
"id":
|
143 |
"content": "c'16 ",
|
144 |
"single_word": false,
|
145 |
"lstrip": false,
|
@@ -148,7 +121,7 @@
|
|
148 |
"special": false
|
149 |
},
|
150 |
{
|
151 |
-
"id":
|
152 |
"content": "c''1 ",
|
153 |
"single_word": false,
|
154 |
"lstrip": false,
|
@@ -157,7 +130,7 @@
|
|
157 |
"special": false
|
158 |
},
|
159 |
{
|
160 |
-
"id":
|
161 |
"content": "c''2 ",
|
162 |
"single_word": false,
|
163 |
"lstrip": false,
|
@@ -166,7 +139,7 @@
|
|
166 |
"special": false
|
167 |
},
|
168 |
{
|
169 |
-
"id":
|
170 |
"content": "c''4 ",
|
171 |
"single_word": false,
|
172 |
"lstrip": false,
|
@@ -175,7 +148,7 @@
|
|
175 |
"special": false
|
176 |
},
|
177 |
{
|
178 |
-
"id":
|
179 |
"content": "c''8 ",
|
180 |
"single_word": false,
|
181 |
"lstrip": false,
|
@@ -184,7 +157,7 @@
|
|
184 |
"special": false
|
185 |
},
|
186 |
{
|
187 |
-
"id":
|
188 |
"content": "c''16 ",
|
189 |
"single_word": false,
|
190 |
"lstrip": false,
|
@@ -193,7 +166,7 @@
|
|
193 |
"special": false
|
194 |
},
|
195 |
{
|
196 |
-
"id":
|
197 |
"content": "d'1 ",
|
198 |
"single_word": false,
|
199 |
"lstrip": false,
|
@@ -202,7 +175,7 @@
|
|
202 |
"special": false
|
203 |
},
|
204 |
{
|
205 |
-
"id":
|
206 |
"content": "d'2 ",
|
207 |
"single_word": false,
|
208 |
"lstrip": false,
|
@@ -211,7 +184,7 @@
|
|
211 |
"special": false
|
212 |
},
|
213 |
{
|
214 |
-
"id":
|
215 |
"content": "d'4 ",
|
216 |
"single_word": false,
|
217 |
"lstrip": false,
|
@@ -220,7 +193,7 @@
|
|
220 |
"special": false
|
221 |
},
|
222 |
{
|
223 |
-
"id":
|
224 |
"content": "d'8 ",
|
225 |
"single_word": false,
|
226 |
"lstrip": false,
|
@@ -229,7 +202,7 @@
|
|
229 |
"special": false
|
230 |
},
|
231 |
{
|
232 |
-
"id":
|
233 |
"content": "d'16 ",
|
234 |
"single_word": false,
|
235 |
"lstrip": false,
|
@@ -238,7 +211,7 @@
|
|
238 |
"special": false
|
239 |
},
|
240 |
{
|
241 |
-
"id":
|
242 |
"content": "d''1 ",
|
243 |
"single_word": false,
|
244 |
"lstrip": false,
|
@@ -247,7 +220,7 @@
|
|
247 |
"special": false
|
248 |
},
|
249 |
{
|
250 |
-
"id":
|
251 |
"content": "d''2 ",
|
252 |
"single_word": false,
|
253 |
"lstrip": false,
|
@@ -256,7 +229,7 @@
|
|
256 |
"special": false
|
257 |
},
|
258 |
{
|
259 |
-
"id":
|
260 |
"content": "d''4 ",
|
261 |
"single_word": false,
|
262 |
"lstrip": false,
|
@@ -265,7 +238,7 @@
|
|
265 |
"special": false
|
266 |
},
|
267 |
{
|
268 |
-
"id":
|
269 |
"content": "d''8 ",
|
270 |
"single_word": false,
|
271 |
"lstrip": false,
|
@@ -274,7 +247,7 @@
|
|
274 |
"special": false
|
275 |
},
|
276 |
{
|
277 |
-
"id":
|
278 |
"content": "d''16 ",
|
279 |
"single_word": false,
|
280 |
"lstrip": false,
|
@@ -283,7 +256,7 @@
|
|
283 |
"special": false
|
284 |
},
|
285 |
{
|
286 |
-
"id":
|
287 |
"content": "e'1 ",
|
288 |
"single_word": false,
|
289 |
"lstrip": false,
|
@@ -292,7 +265,7 @@
|
|
292 |
"special": false
|
293 |
},
|
294 |
{
|
295 |
-
"id":
|
296 |
"content": "e'2 ",
|
297 |
"single_word": false,
|
298 |
"lstrip": false,
|
@@ -301,7 +274,7 @@
|
|
301 |
"special": false
|
302 |
},
|
303 |
{
|
304 |
-
"id":
|
305 |
"content": "e'4 ",
|
306 |
"single_word": false,
|
307 |
"lstrip": false,
|
@@ -310,7 +283,7 @@
|
|
310 |
"special": false
|
311 |
},
|
312 |
{
|
313 |
-
"id":
|
314 |
"content": "e'8 ",
|
315 |
"single_word": false,
|
316 |
"lstrip": false,
|
@@ -319,7 +292,7 @@
|
|
319 |
"special": false
|
320 |
},
|
321 |
{
|
322 |
-
"id":
|
323 |
"content": "e'16 ",
|
324 |
"single_word": false,
|
325 |
"lstrip": false,
|
@@ -328,7 +301,7 @@
|
|
328 |
"special": false
|
329 |
},
|
330 |
{
|
331 |
-
"id":
|
332 |
"content": "e''1 ",
|
333 |
"single_word": false,
|
334 |
"lstrip": false,
|
@@ -337,7 +310,7 @@
|
|
337 |
"special": false
|
338 |
},
|
339 |
{
|
340 |
-
"id":
|
341 |
"content": "e''2 ",
|
342 |
"single_word": false,
|
343 |
"lstrip": false,
|
@@ -346,7 +319,7 @@
|
|
346 |
"special": false
|
347 |
},
|
348 |
{
|
349 |
-
"id":
|
350 |
"content": "e''4 ",
|
351 |
"single_word": false,
|
352 |
"lstrip": false,
|
@@ -355,7 +328,7 @@
|
|
355 |
"special": false
|
356 |
},
|
357 |
{
|
358 |
-
"id":
|
359 |
"content": "e''8 ",
|
360 |
"single_word": false,
|
361 |
"lstrip": false,
|
@@ -364,7 +337,7 @@
|
|
364 |
"special": false
|
365 |
},
|
366 |
{
|
367 |
-
"id":
|
368 |
"content": "e''16 ",
|
369 |
"single_word": false,
|
370 |
"lstrip": false,
|
@@ -373,7 +346,7 @@
|
|
373 |
"special": false
|
374 |
},
|
375 |
{
|
376 |
-
"id":
|
377 |
"content": "f'1 ",
|
378 |
"single_word": false,
|
379 |
"lstrip": false,
|
@@ -382,7 +355,7 @@
|
|
382 |
"special": false
|
383 |
},
|
384 |
{
|
385 |
-
"id":
|
386 |
"content": "f'2 ",
|
387 |
"single_word": false,
|
388 |
"lstrip": false,
|
@@ -391,7 +364,7 @@
|
|
391 |
"special": false
|
392 |
},
|
393 |
{
|
394 |
-
"id":
|
395 |
"content": "f'4 ",
|
396 |
"single_word": false,
|
397 |
"lstrip": false,
|
@@ -400,7 +373,7 @@
|
|
400 |
"special": false
|
401 |
},
|
402 |
{
|
403 |
-
"id":
|
404 |
"content": "f'8 ",
|
405 |
"single_word": false,
|
406 |
"lstrip": false,
|
@@ -409,7 +382,7 @@
|
|
409 |
"special": false
|
410 |
},
|
411 |
{
|
412 |
-
"id":
|
413 |
"content": "f'16 ",
|
414 |
"single_word": false,
|
415 |
"lstrip": false,
|
@@ -418,7 +391,7 @@
|
|
418 |
"special": false
|
419 |
},
|
420 |
{
|
421 |
-
"id":
|
422 |
"content": "f''1 ",
|
423 |
"single_word": false,
|
424 |
"lstrip": false,
|
@@ -427,7 +400,7 @@
|
|
427 |
"special": false
|
428 |
},
|
429 |
{
|
430 |
-
"id":
|
431 |
"content": "f''2 ",
|
432 |
"single_word": false,
|
433 |
"lstrip": false,
|
@@ -436,7 +409,7 @@
|
|
436 |
"special": false
|
437 |
},
|
438 |
{
|
439 |
-
"id":
|
440 |
"content": "f''4 ",
|
441 |
"single_word": false,
|
442 |
"lstrip": false,
|
@@ -445,7 +418,7 @@
|
|
445 |
"special": false
|
446 |
},
|
447 |
{
|
448 |
-
"id":
|
449 |
"content": "f''8 ",
|
450 |
"single_word": false,
|
451 |
"lstrip": false,
|
@@ -454,7 +427,7 @@
|
|
454 |
"special": false
|
455 |
},
|
456 |
{
|
457 |
-
"id":
|
458 |
"content": "f''16 ",
|
459 |
"single_word": false,
|
460 |
"lstrip": false,
|
@@ -463,7 +436,7 @@
|
|
463 |
"special": false
|
464 |
},
|
465 |
{
|
466 |
-
"id":
|
467 |
"content": "g'1 ",
|
468 |
"single_word": false,
|
469 |
"lstrip": false,
|
@@ -472,7 +445,7 @@
|
|
472 |
"special": false
|
473 |
},
|
474 |
{
|
475 |
-
"id":
|
476 |
"content": "g'2 ",
|
477 |
"single_word": false,
|
478 |
"lstrip": false,
|
@@ -481,7 +454,7 @@
|
|
481 |
"special": false
|
482 |
},
|
483 |
{
|
484 |
-
"id":
|
485 |
"content": "g'4 ",
|
486 |
"single_word": false,
|
487 |
"lstrip": false,
|
@@ -490,7 +463,7 @@
|
|
490 |
"special": false
|
491 |
},
|
492 |
{
|
493 |
-
"id":
|
494 |
"content": "g'8 ",
|
495 |
"single_word": false,
|
496 |
"lstrip": false,
|
@@ -499,7 +472,7 @@
|
|
499 |
"special": false
|
500 |
},
|
501 |
{
|
502 |
-
"id":
|
503 |
"content": "g'16 ",
|
504 |
"single_word": false,
|
505 |
"lstrip": false,
|
@@ -508,7 +481,7 @@
|
|
508 |
"special": false
|
509 |
},
|
510 |
{
|
511 |
-
"id":
|
512 |
"content": "g''1 ",
|
513 |
"single_word": false,
|
514 |
"lstrip": false,
|
@@ -517,7 +490,7 @@
|
|
517 |
"special": false
|
518 |
},
|
519 |
{
|
520 |
-
"id":
|
521 |
"content": "g''2 ",
|
522 |
"single_word": false,
|
523 |
"lstrip": false,
|
@@ -526,7 +499,7 @@
|
|
526 |
"special": false
|
527 |
},
|
528 |
{
|
529 |
-
"id":
|
530 |
"content": "g''4 ",
|
531 |
"single_word": false,
|
532 |
"lstrip": false,
|
@@ -535,7 +508,7 @@
|
|
535 |
"special": false
|
536 |
},
|
537 |
{
|
538 |
-
"id":
|
539 |
"content": "g''8 ",
|
540 |
"single_word": false,
|
541 |
"lstrip": false,
|
@@ -544,7 +517,7 @@
|
|
544 |
"special": false
|
545 |
},
|
546 |
{
|
547 |
-
"id":
|
548 |
"content": "g''16 ",
|
549 |
"single_word": false,
|
550 |
"lstrip": false,
|
@@ -553,7 +526,7 @@
|
|
553 |
"special": false
|
554 |
},
|
555 |
{
|
556 |
-
"id":
|
557 |
"content": "a'1 ",
|
558 |
"single_word": false,
|
559 |
"lstrip": false,
|
@@ -562,7 +535,7 @@
|
|
562 |
"special": false
|
563 |
},
|
564 |
{
|
565 |
-
"id":
|
566 |
"content": "a'2 ",
|
567 |
"single_word": false,
|
568 |
"lstrip": false,
|
@@ -571,7 +544,7 @@
|
|
571 |
"special": false
|
572 |
},
|
573 |
{
|
574 |
-
"id":
|
575 |
"content": "a'4 ",
|
576 |
"single_word": false,
|
577 |
"lstrip": false,
|
@@ -580,7 +553,7 @@
|
|
580 |
"special": false
|
581 |
},
|
582 |
{
|
583 |
-
"id":
|
584 |
"content": "a'8 ",
|
585 |
"single_word": false,
|
586 |
"lstrip": false,
|
@@ -589,7 +562,7 @@
|
|
589 |
"special": false
|
590 |
},
|
591 |
{
|
592 |
-
"id":
|
593 |
"content": "a'16 ",
|
594 |
"single_word": false,
|
595 |
"lstrip": false,
|
@@ -598,7 +571,7 @@
|
|
598 |
"special": false
|
599 |
},
|
600 |
{
|
601 |
-
"id":
|
602 |
"content": "a''1 ",
|
603 |
"single_word": false,
|
604 |
"lstrip": false,
|
@@ -607,7 +580,7 @@
|
|
607 |
"special": false
|
608 |
},
|
609 |
{
|
610 |
-
"id":
|
611 |
"content": "a''2 ",
|
612 |
"single_word": false,
|
613 |
"lstrip": false,
|
@@ -616,7 +589,7 @@
|
|
616 |
"special": false
|
617 |
},
|
618 |
{
|
619 |
-
"id":
|
620 |
"content": "a''4 ",
|
621 |
"single_word": false,
|
622 |
"lstrip": false,
|
@@ -625,7 +598,7 @@
|
|
625 |
"special": false
|
626 |
},
|
627 |
{
|
628 |
-
"id":
|
629 |
"content": "a''8 ",
|
630 |
"single_word": false,
|
631 |
"lstrip": false,
|
@@ -634,7 +607,7 @@
|
|
634 |
"special": false
|
635 |
},
|
636 |
{
|
637 |
-
"id":
|
638 |
"content": "a''16 ",
|
639 |
"single_word": false,
|
640 |
"lstrip": false,
|
@@ -643,7 +616,7 @@
|
|
643 |
"special": false
|
644 |
},
|
645 |
{
|
646 |
-
"id":
|
647 |
"content": "b'1 ",
|
648 |
"single_word": false,
|
649 |
"lstrip": false,
|
@@ -652,7 +625,7 @@
|
|
652 |
"special": false
|
653 |
},
|
654 |
{
|
655 |
-
"id":
|
656 |
"content": "b'2 ",
|
657 |
"single_word": false,
|
658 |
"lstrip": false,
|
@@ -661,7 +634,7 @@
|
|
661 |
"special": false
|
662 |
},
|
663 |
{
|
664 |
-
"id":
|
665 |
"content": "b'4 ",
|
666 |
"single_word": false,
|
667 |
"lstrip": false,
|
@@ -670,7 +643,7 @@
|
|
670 |
"special": false
|
671 |
},
|
672 |
{
|
673 |
-
"id":
|
674 |
"content": "b'8 ",
|
675 |
"single_word": false,
|
676 |
"lstrip": false,
|
@@ -679,7 +652,7 @@
|
|
679 |
"special": false
|
680 |
},
|
681 |
{
|
682 |
-
"id":
|
683 |
"content": "b'16 ",
|
684 |
"single_word": false,
|
685 |
"lstrip": false,
|
@@ -688,7 +661,7 @@
|
|
688 |
"special": false
|
689 |
},
|
690 |
{
|
691 |
-
"id":
|
692 |
"content": "b''1 ",
|
693 |
"single_word": false,
|
694 |
"lstrip": false,
|
@@ -697,7 +670,7 @@
|
|
697 |
"special": false
|
698 |
},
|
699 |
{
|
700 |
-
"id":
|
701 |
"content": "b''2 ",
|
702 |
"single_word": false,
|
703 |
"lstrip": false,
|
@@ -706,7 +679,7 @@
|
|
706 |
"special": false
|
707 |
},
|
708 |
{
|
709 |
-
"id":
|
710 |
"content": "b''4 ",
|
711 |
"single_word": false,
|
712 |
"lstrip": false,
|
@@ -715,7 +688,7 @@
|
|
715 |
"special": false
|
716 |
},
|
717 |
{
|
718 |
-
"id":
|
719 |
"content": "b''8 ",
|
720 |
"single_word": false,
|
721 |
"lstrip": false,
|
@@ -724,7 +697,7 @@
|
|
724 |
"special": false
|
725 |
},
|
726 |
{
|
727 |
-
"id":
|
728 |
"content": "b''16 ",
|
729 |
"single_word": false,
|
730 |
"lstrip": false,
|
@@ -733,7 +706,7 @@
|
|
733 |
"special": false
|
734 |
},
|
735 |
{
|
736 |
-
"id":
|
737 |
"content": "r1 ",
|
738 |
"single_word": false,
|
739 |
"lstrip": false,
|
@@ -742,7 +715,7 @@
|
|
742 |
"special": false
|
743 |
},
|
744 |
{
|
745 |
-
"id":
|
746 |
"content": "r2 ",
|
747 |
"single_word": false,
|
748 |
"lstrip": false,
|
@@ -751,7 +724,7 @@
|
|
751 |
"special": false
|
752 |
},
|
753 |
{
|
754 |
-
"id":
|
755 |
"content": "r4 ",
|
756 |
"single_word": false,
|
757 |
"lstrip": false,
|
@@ -760,7 +733,7 @@
|
|
760 |
"special": false
|
761 |
},
|
762 |
{
|
763 |
-
"id":
|
764 |
"content": "r8 ",
|
765 |
"single_word": false,
|
766 |
"lstrip": false,
|
@@ -769,7 +742,7 @@
|
|
769 |
"special": false
|
770 |
},
|
771 |
{
|
772 |
-
"id":
|
773 |
"content": "r16 ",
|
774 |
"single_word": false,
|
775 |
"lstrip": false,
|
|
|
39 |
"normalized": false,
|
40 |
"special": true
|
41 |
},
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
42 |
{
|
43 |
"id": 57521,
|
44 |
"content": "<mask>",
|
|
|
77 |
},
|
78 |
{
|
79 |
"id": 57525,
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
80 |
"content": "c'1 ",
|
81 |
"single_word": false,
|
82 |
"lstrip": false,
|
|
|
85 |
"special": false
|
86 |
},
|
87 |
{
|
88 |
+
"id": 57526,
|
89 |
"content": "c'2 ",
|
90 |
"single_word": false,
|
91 |
"lstrip": false,
|
|
|
94 |
"special": false
|
95 |
},
|
96 |
{
|
97 |
+
"id": 57527,
|
98 |
"content": "c'4 ",
|
99 |
"single_word": false,
|
100 |
"lstrip": false,
|
|
|
103 |
"special": false
|
104 |
},
|
105 |
{
|
106 |
+
"id": 57528,
|
107 |
"content": "c'8 ",
|
108 |
"single_word": false,
|
109 |
"lstrip": false,
|
|
|
112 |
"special": false
|
113 |
},
|
114 |
{
|
115 |
+
"id": 57529,
|
116 |
"content": "c'16 ",
|
117 |
"single_word": false,
|
118 |
"lstrip": false,
|
|
|
121 |
"special": false
|
122 |
},
|
123 |
{
|
124 |
+
"id": 57530,
|
125 |
"content": "c''1 ",
|
126 |
"single_word": false,
|
127 |
"lstrip": false,
|
|
|
130 |
"special": false
|
131 |
},
|
132 |
{
|
133 |
+
"id": 57531,
|
134 |
"content": "c''2 ",
|
135 |
"single_word": false,
|
136 |
"lstrip": false,
|
|
|
139 |
"special": false
|
140 |
},
|
141 |
{
|
142 |
+
"id": 57532,
|
143 |
"content": "c''4 ",
|
144 |
"single_word": false,
|
145 |
"lstrip": false,
|
|
|
148 |
"special": false
|
149 |
},
|
150 |
{
|
151 |
+
"id": 57533,
|
152 |
"content": "c''8 ",
|
153 |
"single_word": false,
|
154 |
"lstrip": false,
|
|
|
157 |
"special": false
|
158 |
},
|
159 |
{
|
160 |
+
"id": 57534,
|
161 |
"content": "c''16 ",
|
162 |
"single_word": false,
|
163 |
"lstrip": false,
|
|
|
166 |
"special": false
|
167 |
},
|
168 |
{
|
169 |
+
"id": 57535,
|
170 |
"content": "d'1 ",
|
171 |
"single_word": false,
|
172 |
"lstrip": false,
|
|
|
175 |
"special": false
|
176 |
},
|
177 |
{
|
178 |
+
"id": 57536,
|
179 |
"content": "d'2 ",
|
180 |
"single_word": false,
|
181 |
"lstrip": false,
|
|
|
184 |
"special": false
|
185 |
},
|
186 |
{
|
187 |
+
"id": 57537,
|
188 |
"content": "d'4 ",
|
189 |
"single_word": false,
|
190 |
"lstrip": false,
|
|
|
193 |
"special": false
|
194 |
},
|
195 |
{
|
196 |
+
"id": 57538,
|
197 |
"content": "d'8 ",
|
198 |
"single_word": false,
|
199 |
"lstrip": false,
|
|
|
202 |
"special": false
|
203 |
},
|
204 |
{
|
205 |
+
"id": 57539,
|
206 |
"content": "d'16 ",
|
207 |
"single_word": false,
|
208 |
"lstrip": false,
|
|
|
211 |
"special": false
|
212 |
},
|
213 |
{
|
214 |
+
"id": 57540,
|
215 |
"content": "d''1 ",
|
216 |
"single_word": false,
|
217 |
"lstrip": false,
|
|
|
220 |
"special": false
|
221 |
},
|
222 |
{
|
223 |
+
"id": 57541,
|
224 |
"content": "d''2 ",
|
225 |
"single_word": false,
|
226 |
"lstrip": false,
|
|
|
229 |
"special": false
|
230 |
},
|
231 |
{
|
232 |
+
"id": 57542,
|
233 |
"content": "d''4 ",
|
234 |
"single_word": false,
|
235 |
"lstrip": false,
|
|
|
238 |
"special": false
|
239 |
},
|
240 |
{
|
241 |
+
"id": 57543,
|
242 |
"content": "d''8 ",
|
243 |
"single_word": false,
|
244 |
"lstrip": false,
|
|
|
247 |
"special": false
|
248 |
},
|
249 |
{
|
250 |
+
"id": 57544,
|
251 |
"content": "d''16 ",
|
252 |
"single_word": false,
|
253 |
"lstrip": false,
|
|
|
256 |
"special": false
|
257 |
},
|
258 |
{
|
259 |
+
"id": 57545,
|
260 |
"content": "e'1 ",
|
261 |
"single_word": false,
|
262 |
"lstrip": false,
|
|
|
265 |
"special": false
|
266 |
},
|
267 |
{
|
268 |
+
"id": 57546,
|
269 |
"content": "e'2 ",
|
270 |
"single_word": false,
|
271 |
"lstrip": false,
|
|
|
274 |
"special": false
|
275 |
},
|
276 |
{
|
277 |
+
"id": 57547,
|
278 |
"content": "e'4 ",
|
279 |
"single_word": false,
|
280 |
"lstrip": false,
|
|
|
283 |
"special": false
|
284 |
},
|
285 |
{
|
286 |
+
"id": 57548,
|
287 |
"content": "e'8 ",
|
288 |
"single_word": false,
|
289 |
"lstrip": false,
|
|
|
292 |
"special": false
|
293 |
},
|
294 |
{
|
295 |
+
"id": 57549,
|
296 |
"content": "e'16 ",
|
297 |
"single_word": false,
|
298 |
"lstrip": false,
|
|
|
301 |
"special": false
|
302 |
},
|
303 |
{
|
304 |
+
"id": 57550,
|
305 |
"content": "e''1 ",
|
306 |
"single_word": false,
|
307 |
"lstrip": false,
|
|
|
310 |
"special": false
|
311 |
},
|
312 |
{
|
313 |
+
"id": 57551,
|
314 |
"content": "e''2 ",
|
315 |
"single_word": false,
|
316 |
"lstrip": false,
|
|
|
319 |
"special": false
|
320 |
},
|
321 |
{
|
322 |
+
"id": 57552,
|
323 |
"content": "e''4 ",
|
324 |
"single_word": false,
|
325 |
"lstrip": false,
|
|
|
328 |
"special": false
|
329 |
},
|
330 |
{
|
331 |
+
"id": 57553,
|
332 |
"content": "e''8 ",
|
333 |
"single_word": false,
|
334 |
"lstrip": false,
|
|
|
337 |
"special": false
|
338 |
},
|
339 |
{
|
340 |
+
"id": 57554,
|
341 |
"content": "e''16 ",
|
342 |
"single_word": false,
|
343 |
"lstrip": false,
|
|
|
346 |
"special": false
|
347 |
},
|
348 |
{
|
349 |
+
"id": 57555,
|
350 |
"content": "f'1 ",
|
351 |
"single_word": false,
|
352 |
"lstrip": false,
|
|
|
355 |
"special": false
|
356 |
},
|
357 |
{
|
358 |
+
"id": 57556,
|
359 |
"content": "f'2 ",
|
360 |
"single_word": false,
|
361 |
"lstrip": false,
|
|
|
364 |
"special": false
|
365 |
},
|
366 |
{
|
367 |
+
"id": 57557,
|
368 |
"content": "f'4 ",
|
369 |
"single_word": false,
|
370 |
"lstrip": false,
|
|
|
373 |
"special": false
|
374 |
},
|
375 |
{
|
376 |
+
"id": 57558,
|
377 |
"content": "f'8 ",
|
378 |
"single_word": false,
|
379 |
"lstrip": false,
|
|
|
382 |
"special": false
|
383 |
},
|
384 |
{
|
385 |
+
"id": 57559,
|
386 |
"content": "f'16 ",
|
387 |
"single_word": false,
|
388 |
"lstrip": false,
|
|
|
391 |
"special": false
|
392 |
},
|
393 |
{
|
394 |
+
"id": 57560,
|
395 |
"content": "f''1 ",
|
396 |
"single_word": false,
|
397 |
"lstrip": false,
|
|
|
400 |
"special": false
|
401 |
},
|
402 |
{
|
403 |
+
"id": 57561,
|
404 |
"content": "f''2 ",
|
405 |
"single_word": false,
|
406 |
"lstrip": false,
|
|
|
409 |
"special": false
|
410 |
},
|
411 |
{
|
412 |
+
"id": 57562,
|
413 |
"content": "f''4 ",
|
414 |
"single_word": false,
|
415 |
"lstrip": false,
|
|
|
418 |
"special": false
|
419 |
},
|
420 |
{
|
421 |
+
"id": 57563,
|
422 |
"content": "f''8 ",
|
423 |
"single_word": false,
|
424 |
"lstrip": false,
|
|
|
427 |
"special": false
|
428 |
},
|
429 |
{
|
430 |
+
"id": 57564,
|
431 |
"content": "f''16 ",
|
432 |
"single_word": false,
|
433 |
"lstrip": false,
|
|
|
436 |
"special": false
|
437 |
},
|
438 |
{
|
439 |
+
"id": 57565,
|
440 |
"content": "g'1 ",
|
441 |
"single_word": false,
|
442 |
"lstrip": false,
|
|
|
445 |
"special": false
|
446 |
},
|
447 |
{
|
448 |
+
"id": 57566,
|
449 |
"content": "g'2 ",
|
450 |
"single_word": false,
|
451 |
"lstrip": false,
|
|
|
454 |
"special": false
|
455 |
},
|
456 |
{
|
457 |
+
"id": 57567,
|
458 |
"content": "g'4 ",
|
459 |
"single_word": false,
|
460 |
"lstrip": false,
|
|
|
463 |
"special": false
|
464 |
},
|
465 |
{
|
466 |
+
"id": 57568,
|
467 |
"content": "g'8 ",
|
468 |
"single_word": false,
|
469 |
"lstrip": false,
|
|
|
472 |
"special": false
|
473 |
},
|
474 |
{
|
475 |
+
"id": 57569,
|
476 |
"content": "g'16 ",
|
477 |
"single_word": false,
|
478 |
"lstrip": false,
|
|
|
481 |
"special": false
|
482 |
},
|
483 |
{
|
484 |
+
"id": 57570,
|
485 |
"content": "g''1 ",
|
486 |
"single_word": false,
|
487 |
"lstrip": false,
|
|
|
490 |
"special": false
|
491 |
},
|
492 |
{
|
493 |
+
"id": 57571,
|
494 |
"content": "g''2 ",
|
495 |
"single_word": false,
|
496 |
"lstrip": false,
|
|
|
499 |
"special": false
|
500 |
},
|
501 |
{
|
502 |
+
"id": 57572,
|
503 |
"content": "g''4 ",
|
504 |
"single_word": false,
|
505 |
"lstrip": false,
|
|
|
508 |
"special": false
|
509 |
},
|
510 |
{
|
511 |
+
"id": 57573,
|
512 |
"content": "g''8 ",
|
513 |
"single_word": false,
|
514 |
"lstrip": false,
|
|
|
517 |
"special": false
|
518 |
},
|
519 |
{
|
520 |
+
"id": 57574,
|
521 |
"content": "g''16 ",
|
522 |
"single_word": false,
|
523 |
"lstrip": false,
|
|
|
526 |
"special": false
|
527 |
},
|
528 |
{
|
529 |
+
"id": 57575,
|
530 |
"content": "a'1 ",
|
531 |
"single_word": false,
|
532 |
"lstrip": false,
|
|
|
535 |
"special": false
|
536 |
},
|
537 |
{
|
538 |
+
"id": 57576,
|
539 |
"content": "a'2 ",
|
540 |
"single_word": false,
|
541 |
"lstrip": false,
|
|
|
544 |
"special": false
|
545 |
},
|
546 |
{
|
547 |
+
"id": 57577,
|
548 |
"content": "a'4 ",
|
549 |
"single_word": false,
|
550 |
"lstrip": false,
|
|
|
553 |
"special": false
|
554 |
},
|
555 |
{
|
556 |
+
"id": 57578,
|
557 |
"content": "a'8 ",
|
558 |
"single_word": false,
|
559 |
"lstrip": false,
|
|
|
562 |
"special": false
|
563 |
},
|
564 |
{
|
565 |
+
"id": 57579,
|
566 |
"content": "a'16 ",
|
567 |
"single_word": false,
|
568 |
"lstrip": false,
|
|
|
571 |
"special": false
|
572 |
},
|
573 |
{
|
574 |
+
"id": 57580,
|
575 |
"content": "a''1 ",
|
576 |
"single_word": false,
|
577 |
"lstrip": false,
|
|
|
580 |
"special": false
|
581 |
},
|
582 |
{
|
583 |
+
"id": 57581,
|
584 |
"content": "a''2 ",
|
585 |
"single_word": false,
|
586 |
"lstrip": false,
|
|
|
589 |
"special": false
|
590 |
},
|
591 |
{
|
592 |
+
"id": 57582,
|
593 |
"content": "a''4 ",
|
594 |
"single_word": false,
|
595 |
"lstrip": false,
|
|
|
598 |
"special": false
|
599 |
},
|
600 |
{
|
601 |
+
"id": 57583,
|
602 |
"content": "a''8 ",
|
603 |
"single_word": false,
|
604 |
"lstrip": false,
|
|
|
607 |
"special": false
|
608 |
},
|
609 |
{
|
610 |
+
"id": 57584,
|
611 |
"content": "a''16 ",
|
612 |
"single_word": false,
|
613 |
"lstrip": false,
|
|
|
616 |
"special": false
|
617 |
},
|
618 |
{
|
619 |
+
"id": 57585,
|
620 |
"content": "b'1 ",
|
621 |
"single_word": false,
|
622 |
"lstrip": false,
|
|
|
625 |
"special": false
|
626 |
},
|
627 |
{
|
628 |
+
"id": 57586,
|
629 |
"content": "b'2 ",
|
630 |
"single_word": false,
|
631 |
"lstrip": false,
|
|
|
634 |
"special": false
|
635 |
},
|
636 |
{
|
637 |
+
"id": 57587,
|
638 |
"content": "b'4 ",
|
639 |
"single_word": false,
|
640 |
"lstrip": false,
|
|
|
643 |
"special": false
|
644 |
},
|
645 |
{
|
646 |
+
"id": 57588,
|
647 |
"content": "b'8 ",
|
648 |
"single_word": false,
|
649 |
"lstrip": false,
|
|
|
652 |
"special": false
|
653 |
},
|
654 |
{
|
655 |
+
"id": 57589,
|
656 |
"content": "b'16 ",
|
657 |
"single_word": false,
|
658 |
"lstrip": false,
|
|
|
661 |
"special": false
|
662 |
},
|
663 |
{
|
664 |
+
"id": 57590,
|
665 |
"content": "b''1 ",
|
666 |
"single_word": false,
|
667 |
"lstrip": false,
|
|
|
670 |
"special": false
|
671 |
},
|
672 |
{
|
673 |
+
"id": 57591,
|
674 |
"content": "b''2 ",
|
675 |
"single_word": false,
|
676 |
"lstrip": false,
|
|
|
679 |
"special": false
|
680 |
},
|
681 |
{
|
682 |
+
"id": 57592,
|
683 |
"content": "b''4 ",
|
684 |
"single_word": false,
|
685 |
"lstrip": false,
|
|
|
688 |
"special": false
|
689 |
},
|
690 |
{
|
691 |
+
"id": 57593,
|
692 |
"content": "b''8 ",
|
693 |
"single_word": false,
|
694 |
"lstrip": false,
|
|
|
697 |
"special": false
|
698 |
},
|
699 |
{
|
700 |
+
"id": 57594,
|
701 |
"content": "b''16 ",
|
702 |
"single_word": false,
|
703 |
"lstrip": false,
|
|
|
706 |
"special": false
|
707 |
},
|
708 |
{
|
709 |
+
"id": 57595,
|
710 |
"content": "r1 ",
|
711 |
"single_word": false,
|
712 |
"lstrip": false,
|
|
|
715 |
"special": false
|
716 |
},
|
717 |
{
|
718 |
+
"id": 57596,
|
719 |
"content": "r2 ",
|
720 |
"single_word": false,
|
721 |
"lstrip": false,
|
|
|
724 |
"special": false
|
725 |
},
|
726 |
{
|
727 |
+
"id": 57597,
|
728 |
"content": "r4 ",
|
729 |
"single_word": false,
|
730 |
"lstrip": false,
|
|
|
733 |
"special": false
|
734 |
},
|
735 |
{
|
736 |
+
"id": 57598,
|
737 |
"content": "r8 ",
|
738 |
"single_word": false,
|
739 |
"lstrip": false,
|
|
|
742 |
"special": false
|
743 |
},
|
744 |
{
|
745 |
+
"id": 57599,
|
746 |
"content": "r16 ",
|
747 |
"single_word": false,
|
748 |
"lstrip": false,
|
tokenizer_config.json
CHANGED
@@ -11,7 +11,7 @@
|
|
11 |
"single_word": false
|
12 |
},
|
13 |
"model_max_length": 1000000000000000019884624838656,
|
14 |
-
"name_or_path": "./
|
15 |
"pad_token": "<pad>",
|
16 |
"processor_class": "DonutProcessor",
|
17 |
"sep_token": "</s>",
|
|
|
11 |
"single_word": false
|
12 |
},
|
13 |
"model_max_length": 1000000000000000019884624838656,
|
14 |
+
"name_or_path": "./model_7",
|
15 |
"pad_token": "<pad>",
|
16 |
"processor_class": "DonutProcessor",
|
17 |
"sep_token": "</s>",
|