Commit
·
a1850fc
1
Parent(s):
7cc7144
Upload tokenizer
Browse files- added_tokens.json +5 -0
- tokenizer.json +28 -6
added_tokens.json
ADDED
@@ -0,0 +1,5 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"p0": 50259,
|
3 |
+
"p1": 50257,
|
4 |
+
"p2": 50258
|
5 |
+
}
|
tokenizer.json
CHANGED
@@ -1,11 +1,6 @@
|
|
1 |
{
|
2 |
"version": "1.0",
|
3 |
-
"truncation":
|
4 |
-
"direction": "Right",
|
5 |
-
"max_length": 1024,
|
6 |
-
"strategy": "LongestFirst",
|
7 |
-
"stride": 0
|
8 |
-
},
|
9 |
"padding": null,
|
10 |
"added_tokens": [
|
11 |
{
|
@@ -16,6 +11,33 @@
|
|
16 |
"rstrip": false,
|
17 |
"normalized": false,
|
18 |
"special": true
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
19 |
}
|
20 |
],
|
21 |
"normalizer": null,
|
|
|
1 |
{
|
2 |
"version": "1.0",
|
3 |
+
"truncation": null,
|
|
|
|
|
|
|
|
|
|
|
4 |
"padding": null,
|
5 |
"added_tokens": [
|
6 |
{
|
|
|
11 |
"rstrip": false,
|
12 |
"normalized": false,
|
13 |
"special": true
|
14 |
+
},
|
15 |
+
{
|
16 |
+
"id": 50257,
|
17 |
+
"content": "p1",
|
18 |
+
"single_word": false,
|
19 |
+
"lstrip": false,
|
20 |
+
"rstrip": false,
|
21 |
+
"normalized": true,
|
22 |
+
"special": false
|
23 |
+
},
|
24 |
+
{
|
25 |
+
"id": 50258,
|
26 |
+
"content": "p2",
|
27 |
+
"single_word": false,
|
28 |
+
"lstrip": false,
|
29 |
+
"rstrip": false,
|
30 |
+
"normalized": true,
|
31 |
+
"special": false
|
32 |
+
},
|
33 |
+
{
|
34 |
+
"id": 50259,
|
35 |
+
"content": "p0",
|
36 |
+
"single_word": false,
|
37 |
+
"lstrip": false,
|
38 |
+
"rstrip": false,
|
39 |
+
"normalized": true,
|
40 |
+
"special": false
|
41 |
}
|
42 |
],
|
43 |
"normalizer": null,
|