TheRains commited on
Commit
2de08a8
1 Parent(s): 614d785

End of training

Browse files
added_tokens.json CHANGED
@@ -17,7 +17,6 @@
17
  "<|da|>": 50285,
18
  "<|de|>": 50261,
19
  "<|el|>": 50281,
20
- "<|endoftext|>": 50257,
21
  "<|en|>": 50259,
22
  "<|es|>": 50262,
23
  "<|et|>": 50307,
@@ -30,6 +29,7 @@
30
  "<|gu|>": 50333,
31
  "<|haw|>": 50352,
32
  "<|ha|>": 50354,
 
33
  "<|hi|>": 50276,
34
  "<|hr|>": 50291,
35
  "<|ht|>": 50339,
@@ -38,7 +38,6 @@
38
  "<|id|>": 50275,
39
  "<|is|>": 50311,
40
  "<|it|>": 50274,
41
- "<|iw|>": 50279,
42
  "<|ja|>": 50266,
43
  "<|jw|>": 50356,
44
  "<|ka|>": 50329,
 
17
  "<|da|>": 50285,
18
  "<|de|>": 50261,
19
  "<|el|>": 50281,
 
20
  "<|en|>": 50259,
21
  "<|es|>": 50262,
22
  "<|et|>": 50307,
 
29
  "<|gu|>": 50333,
30
  "<|haw|>": 50352,
31
  "<|ha|>": 50354,
32
+ "<|he|>": 50279,
33
  "<|hi|>": 50276,
34
  "<|hr|>": 50291,
35
  "<|ht|>": 50339,
 
38
  "<|id|>": 50275,
39
  "<|is|>": 50311,
40
  "<|it|>": 50274,
 
41
  "<|ja|>": 50266,
42
  "<|jw|>": 50356,
43
  "<|ka|>": 50329,
generation_config.json CHANGED
@@ -1,12 +1,11 @@
1
  {
2
- "_from_model_config": true,
3
  "begin_suppress_tokens": [
4
  220,
5
- 50256
6
  ],
7
  "bos_token_id": 50257,
8
- "decoder_start_token_id": 50257,
9
- "eos_token_id": 50256,
10
  "forced_decoder_ids": [
11
  [
12
  1,
@@ -14,11 +13,116 @@
14
  ],
15
  [
16
  2,
17
- 50358
18
  ]
19
  ],
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
20
  "max_length": 448,
 
21
  "pad_token_id": 50256,
 
22
  "suppress_tokens": [
23
  1,
24
  2,
@@ -43,71 +147,73 @@
43
  91,
44
  92,
45
  93,
46
- 357,
47
- 366,
48
- 438,
49
- 532,
50
- 685,
51
- 705,
52
- 796,
53
- 930,
54
- 1058,
55
- 1220,
56
- 1267,
57
- 1279,
58
- 1303,
59
- 1343,
60
- 1377,
61
- 1391,
62
- 1635,
63
- 1782,
64
- 1875,
65
- 2162,
66
- 2361,
67
- 2488,
68
- 3467,
69
- 4008,
70
- 4211,
71
- 4600,
72
- 4808,
73
- 5299,
74
- 5855,
75
- 6329,
76
- 7203,
77
- 9609,
78
- 9959,
79
- 10563,
80
- 10786,
81
- 11420,
82
- 11709,
83
- 11907,
84
- 13163,
85
- 13697,
86
- 13700,
87
- 14808,
88
- 15306,
89
- 16410,
90
- 16791,
91
- 17992,
92
- 19203,
93
- 19510,
94
- 20724,
95
- 22305,
96
- 22935,
97
- 27007,
98
- 30109,
99
- 30420,
100
- 33409,
101
- 34949,
102
- 40283,
103
- 40493,
104
- 40549,
105
- 47282,
106
- 49146,
107
- 50257,
108
- 50359,
109
  50360,
110
- 50361
 
111
  ],
 
 
 
 
112
  "transformers_version": "4.27.0.dev0"
113
  }
 
1
  {
 
2
  "begin_suppress_tokens": [
3
  220,
4
+ 50257
5
  ],
6
  "bos_token_id": 50257,
7
+ "decoder_start_token_id": 50258,
8
+ "eos_token_id": 50257,
9
  "forced_decoder_ids": [
10
  [
11
  1,
 
13
  ],
14
  [
15
  2,
16
+ 50359
17
  ]
18
  ],
19
+ "is_multilingual": true,
20
+ "lang_to_id": {
21
+ "<|af|>": 50327,
22
+ "<|am|>": 50334,
23
+ "<|ar|>": 50272,
24
+ "<|as|>": 50350,
25
+ "<|az|>": 50304,
26
+ "<|ba|>": 50355,
27
+ "<|be|>": 50330,
28
+ "<|bg|>": 50292,
29
+ "<|bn|>": 50302,
30
+ "<|bo|>": 50347,
31
+ "<|br|>": 50309,
32
+ "<|bs|>": 50315,
33
+ "<|ca|>": 50270,
34
+ "<|cs|>": 50283,
35
+ "<|cy|>": 50297,
36
+ "<|da|>": 50285,
37
+ "<|de|>": 50261,
38
+ "<|el|>": 50281,
39
+ "<|en|>": 50259,
40
+ "<|es|>": 50262,
41
+ "<|et|>": 50307,
42
+ "<|eu|>": 50310,
43
+ "<|fa|>": 50300,
44
+ "<|fi|>": 50277,
45
+ "<|fo|>": 50338,
46
+ "<|fr|>": 50265,
47
+ "<|gl|>": 50319,
48
+ "<|gu|>": 50333,
49
+ "<|haw|>": 50352,
50
+ "<|ha|>": 50354,
51
+ "<|he|>": 50279,
52
+ "<|hi|>": 50276,
53
+ "<|hr|>": 50291,
54
+ "<|ht|>": 50339,
55
+ "<|hu|>": 50286,
56
+ "<|hy|>": 50312,
57
+ "<|id|>": 50275,
58
+ "<|is|>": 50311,
59
+ "<|it|>": 50274,
60
+ "<|ja|>": 50266,
61
+ "<|jw|>": 50356,
62
+ "<|ka|>": 50329,
63
+ "<|kk|>": 50316,
64
+ "<|km|>": 50323,
65
+ "<|kn|>": 50306,
66
+ "<|ko|>": 50264,
67
+ "<|la|>": 50294,
68
+ "<|lb|>": 50345,
69
+ "<|ln|>": 50353,
70
+ "<|lo|>": 50336,
71
+ "<|lt|>": 50293,
72
+ "<|lv|>": 50301,
73
+ "<|mg|>": 50349,
74
+ "<|mi|>": 50295,
75
+ "<|mk|>": 50308,
76
+ "<|ml|>": 50296,
77
+ "<|mn|>": 50314,
78
+ "<|mr|>": 50320,
79
+ "<|ms|>": 50282,
80
+ "<|mt|>": 50343,
81
+ "<|my|>": 50346,
82
+ "<|ne|>": 50313,
83
+ "<|nl|>": 50271,
84
+ "<|nn|>": 50342,
85
+ "<|no|>": 50288,
86
+ "<|oc|>": 50328,
87
+ "<|pa|>": 50321,
88
+ "<|pl|>": 50269,
89
+ "<|ps|>": 50340,
90
+ "<|pt|>": 50267,
91
+ "<|ro|>": 50284,
92
+ "<|ru|>": 50263,
93
+ "<|sa|>": 50344,
94
+ "<|sd|>": 50332,
95
+ "<|si|>": 50322,
96
+ "<|sk|>": 50298,
97
+ "<|sl|>": 50305,
98
+ "<|sn|>": 50324,
99
+ "<|so|>": 50326,
100
+ "<|sq|>": 50317,
101
+ "<|sr|>": 50303,
102
+ "<|su|>": 50357,
103
+ "<|sv|>": 50273,
104
+ "<|sw|>": 50318,
105
+ "<|ta|>": 50287,
106
+ "<|te|>": 50299,
107
+ "<|tg|>": 50331,
108
+ "<|th|>": 50289,
109
+ "<|tk|>": 50341,
110
+ "<|tl|>": 50348,
111
+ "<|tr|>": 50268,
112
+ "<|tt|>": 50351,
113
+ "<|uk|>": 50280,
114
+ "<|ur|>": 50290,
115
+ "<|uz|>": 50337,
116
+ "<|vi|>": 50278,
117
+ "<|yi|>": 50335,
118
+ "<|yo|>": 50325,
119
+ "<|zh|>": 50260
120
+ },
121
+ "max_initial_timestamp_index": 1,
122
  "max_length": 448,
123
+ "no_timestamps_token_id": 50363,
124
  "pad_token_id": 50256,
125
+ "return_timestamps": false,
126
  "suppress_tokens": [
127
  1,
128
  2,
 
147
  91,
148
  92,
149
  93,
150
+ 359,
151
+ 503,
152
+ 522,
153
+ 542,
154
+ 873,
155
+ 893,
156
+ 902,
157
+ 918,
158
+ 922,
159
+ 931,
160
+ 1350,
161
+ 1853,
162
+ 1982,
163
+ 2460,
164
+ 2627,
165
+ 3246,
166
+ 3253,
167
+ 3268,
168
+ 3536,
169
+ 3846,
170
+ 3961,
171
+ 4183,
172
+ 4667,
173
+ 6585,
174
+ 6647,
175
+ 7273,
176
+ 9061,
177
+ 9383,
178
+ 10428,
179
+ 10929,
180
+ 11938,
181
+ 12033,
182
+ 12331,
183
+ 12562,
184
+ 13793,
185
+ 14157,
186
+ 14635,
187
+ 15265,
188
+ 15618,
189
+ 16553,
190
+ 16604,
191
+ 18362,
192
+ 18956,
193
+ 20075,
194
+ 21675,
195
+ 22520,
196
+ 26130,
197
+ 26161,
198
+ 26435,
199
+ 28279,
200
+ 29464,
201
+ 31650,
202
+ 32302,
203
+ 32470,
204
+ 36865,
205
+ 42863,
206
+ 47425,
207
+ 49870,
208
+ 50254,
209
+ 50258,
 
 
 
210
  50360,
211
+ 50361,
212
+ 50362
213
  ],
214
+ "task_to_id": {
215
+ "transcribe": 50359,
216
+ "translate": 50358
217
+ },
218
  "transformers_version": "4.27.0.dev0"
219
  }
preprocessor_config.json CHANGED
@@ -5,7 +5,7 @@
5
  "hop_length": 160,
6
  "mel_filters": [
7
  [
8
- 0.0,
9
  0.02486259490251541,
10
  0.0,
11
  0.0,
 
5
  "hop_length": 160,
6
  "mel_filters": [
7
  [
8
+ -0.0,
9
  0.02486259490251541,
10
  0.0,
11
  0.0,
pytorch_model.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:3fae540b7be07007361ef03c1129dd6f15f53fb6672b998483205fe9349e69ad
3
  size 967102601
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a2bdd50cba982d9d16b525e43c1f342ff368045ea2e3442f9406e34d22ae73b0
3
  size 967102601
runs/Feb06_17-41-52_60564a25f456/1675705562.9045053/events.out.tfevents.1675705562.60564a25f456.457.1 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:6523c9155373dd4bf73ce8dcf065ccb3db89510b07c37550c3e16f432013239c
3
+ size 5973
runs/Feb06_17-41-52_60564a25f456/events.out.tfevents.1675705562.60564a25f456.457.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e2ffc776c1f5eed8f60da3dff0d3be352206b0e084d1ecc5987ca2f238fed5b6
3
+ size 7236
special_tokens_map.json CHANGED
@@ -22,7 +22,7 @@
22
  "<|hi|>",
23
  "<|fi|>",
24
  "<|vi|>",
25
- "<|iw|>",
26
  "<|uk|>",
27
  "<|el|>",
28
  "<|ms|>",
@@ -124,7 +124,7 @@
124
  },
125
  "pad_token": "<|endoftext|>",
126
  "unk_token": {
127
- "content": "",
128
  "lstrip": false,
129
  "normalized": true,
130
  "rstrip": false,
 
22
  "<|hi|>",
23
  "<|fi|>",
24
  "<|vi|>",
25
+ "<|he|>",
26
  "<|uk|>",
27
  "<|el|>",
28
  "<|ms|>",
 
124
  },
125
  "pad_token": "<|endoftext|>",
126
  "unk_token": {
127
+ "content": "<|endoftext|>",
128
  "lstrip": false,
129
  "normalized": true,
130
  "rstrip": false,
tokenizer_config.json CHANGED
@@ -27,7 +27,7 @@
27
  "tokenizer_class": "WhisperTokenizer",
28
  "unk_token": {
29
  "__type": "AddedToken",
30
- "content": "",
31
  "lstrip": false,
32
  "normalized": true,
33
  "rstrip": false,
 
27
  "tokenizer_class": "WhisperTokenizer",
28
  "unk_token": {
29
  "__type": "AddedToken",
30
+ "content": "<|endoftext|>",
31
  "lstrip": false,
32
  "normalized": true,
33
  "rstrip": false,
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:3f83a44417607667b958d5f2b0a1fff0accfe0cd6ad183593f7a18788a1171c0
3
  size 3643
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:694b8345a6221e59503dd9f7423c13d458fb30c96683dcd1cf296d731f9901a7
3
  size 3643
vocab.json CHANGED
@@ -314,6 +314,7 @@
314
  ";;": 35746,
315
  "<": 27,
316
  "</": 3433,
 
317
  "=": 28,
318
  "=\"": 13114,
319
  "=\"#": 34106,
 
314
  ";;": 35746,
315
  "<": 27,
316
  "</": 3433,
317
+ "<|endoftext|>": 50257,
318
  "=": 28,
319
  "=\"": 13114,
320
  "=\"#": 34106,