ArtusDev commited on
Commit
2b3bfc2
·
verified ·
1 Parent(s): beac6fc

Upload folder using huggingface_hub

Browse files
.gitattributes CHANGED
@@ -33,3 +33,4 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
 
 
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
36
+ tokenizer.json filter=lfs diff=lfs merge=lfs -text
README.md ADDED
@@ -0,0 +1,971 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ ---
2
+ base_model:
3
+ - zerofata/MS3.2-PaintedFantasy-v2-24B
4
+ library_name: transformers
5
+ tags:
6
+ - mergekit
7
+ - merge
8
+ - axolotl
9
+ license: apache-2.0
10
+ datasets:
11
+ - zerofata/Roleplay-Anime-Characters
12
+ - zerofata/Instruct-Anime-CreativeWriting
13
+ - zerofata/Instruct-Anime
14
+ - zerofata/Summaries-Anime-FandomPages
15
+ - zerofata/Stories-Anime
16
+ ---
17
+ <style>
18
+ .container {
19
+ --primary-accent: #EC83B1;
20
+ --secondary-accent: #86C5E5;
21
+ --tertiary-accent: #FDE484;
22
+ --accent-rose: #F8A5C2;
23
+
24
+ --bg-main: #1A1D2E;
25
+ --bg-container: #232741;
26
+ --bg-card: rgba(40, 45, 70, 0.7);
27
+
28
+ --text-main: #E8ECF0;
29
+ --text-muted: #B8C2D0;
30
+ --white: #FFFFFF;
31
+
32
+ --font-title: 'Inter', serif;
33
+ --font-heading: 'Inter', serif;
34
+ --font-body: 'Inter', serif;
35
+ --font-code: 'JetBrains Mono', monospace;
36
+
37
+ font-family: var(--font-body);
38
+ color: var(--text-main);
39
+ line-height: 1.6;
40
+
41
+ max-width: 1200px;
42
+ margin: 20px auto;
43
+ padding: 40px 20px;
44
+ background-color: var(--bg-container);
45
+ background-image:
46
+ radial-gradient(circle at 20% 80%, rgba(236, 131, 177, 0.04) 0%, transparent 50%),
47
+ radial-gradient(circle at 80% 20%, rgba(134, 197, 229, 0.04) 0%, transparent 50%),
48
+ radial-gradient(circle at 40% 40%, rgba(253, 228, 132, 0.02) 0%, transparent 50%);
49
+ min-height: calc(100vh - 40px);
50
+ border: 1px solid var(--primary-accent);
51
+ border-radius: 8px;
52
+ box-shadow: 0 8px 32px rgba(236, 131, 177, 0.07);
53
+ }
54
+
55
+ .container .title-container {
56
+ background-color: var(--bg-main);
57
+ position: relative;
58
+ overflow: hidden;
59
+ margin-bottom: 40px;
60
+ border-left: 3px solid var(--primary-accent);
61
+ box-shadow: 0 6px 20px rgba(236, 131, 177, 0.07);
62
+ }
63
+
64
+ .container .title-wrapper {
65
+ position: relative;
66
+ z-index: 2;
67
+ padding: 25px 20px 30px 30px;
68
+ font-family: var(--font-title);
69
+ }
70
+
71
+ .container .title-main {
72
+ color: var(--accent-rose);
73
+ font-size: 2.5rem;
74
+ font-weight: 700;
75
+ margin: 0;
76
+ letter-spacing: 2px;
77
+ display: inline-block;
78
+ position: relative;
79
+ text-transform: uppercase;
80
+ }
81
+
82
+ .container .title-prefix {
83
+ position: relative;
84
+ z-index: 2;
85
+ }
86
+
87
+ .container .lemonade-text {
88
+ color: var(--secondary-accent);
89
+ position: relative;
90
+ z-index: 2;
91
+ margin-left: 0.2em;
92
+ text-shadow: 0 0 15px var(--secondary-accent);
93
+ }
94
+
95
+ .container .title-subtitle {
96
+ padding-left: 15px;
97
+ margin-top: 5px;
98
+ margin-left: 5px;
99
+ }
100
+
101
+ .container .subtitle-text {
102
+ color: var(--text-muted);
103
+ font-size: 1.2rem;
104
+ font-family: var(--font-body);
105
+ font-weight: 300;
106
+ letter-spacing: 3px;
107
+ text-transform: uppercase;
108
+ display: inline-block;
109
+ }
110
+
111
+ .container .glitchy-overlay {
112
+ position: absolute;
113
+ top: 0;
114
+ left: 0;
115
+ width: 100%;
116
+ height: 100%;
117
+ background-image: repeating-linear-gradient(0deg, rgba(0,0,0,0) 0, rgba(134, 197, 229, 0.08) 1px, rgba(0,0,0,0) 2px);
118
+ z-index: 1;
119
+ }
120
+
121
+ .container img {
122
+ max-width: 100%;
123
+ border: 3px solid var(--white);
124
+ margin-bottom: 30px;
125
+ box-shadow: 0 0 15px rgba(0, 0, 0, 0.3);
126
+ }
127
+
128
+ .container .section-container {
129
+ background-color: var(--bg-card);
130
+ margin-bottom: 30px;
131
+ position: relative;
132
+ overflow: hidden;
133
+ border-bottom: none !important;
134
+ box-shadow: 0 4px 15px rgba(236, 131, 177, 0.05);
135
+ }
136
+
137
+ .container .section-header {
138
+ display: flex;
139
+ align-items: center;
140
+ background-color: rgba(236, 131, 177, 0.12);
141
+ padding: 10px 20px;
142
+ border-bottom: none !important;
143
+ }
144
+
145
+ .container .section-indicator {
146
+ width: 8px;
147
+ height: 20px;
148
+ background-color: var(--primary-accent);
149
+ margin-right: 15px;
150
+ box-shadow: 0 0 8px rgba(236, 131, 177, 0.2);
151
+ }
152
+
153
+ .container .section-title {
154
+ font-family: var(--font-heading);
155
+ color: var(--accent-rose);
156
+ font-size: 1.4rem;
157
+ margin: 0 !important;
158
+ padding: 0 !important;
159
+ letter-spacing: 1px;
160
+ font-weight: 400;
161
+ text-transform: capitalize;
162
+ border-bottom: none !important;
163
+ }
164
+
165
+ .container .section-content {
166
+ padding: 20px;
167
+ font-family: var(--font-body);
168
+ color: var(--text-main);
169
+ line-height: 1.6;
170
+ }
171
+
172
+ .container .subheading {
173
+ color: var(--text-muted);
174
+ font-size: 1.1rem;
175
+ margin-top: 20px;
176
+ margin-bottom: 15px;
177
+ font-weight: 400;
178
+ border-bottom: 1px dashed rgba(184, 194, 208, 0.4);
179
+ display: inline-block;
180
+ text-transform: uppercase;
181
+ letter-spacing: 1px;
182
+ font-family: var(--font-heading);
183
+ }
184
+
185
+ .container .data-box {
186
+ background-color: rgba(26, 29, 46, 0.6);
187
+ padding: 15px;
188
+ border-left: 2px solid var(--primary-accent);
189
+ margin-bottom: 20px;
190
+ box-shadow: 0 2px 10px rgba(236, 131, 177, 0.05);
191
+ }
192
+
193
+ .container .data-row {
194
+ display: flex;
195
+ margin-bottom: 8px;
196
+ align-items: center;
197
+ }
198
+ .container .data-row:last-child { margin-bottom: 0; }
199
+
200
+ .container .data-arrow {
201
+ color: var(--primary-accent);
202
+ width: 20px;
203
+ display: inline-block;
204
+ }
205
+
206
+ .container .data-label {
207
+ color: var(--text-muted);
208
+ width: 80px;
209
+ display: inline-block;
210
+ }
211
+
212
+ .container a {
213
+ color: var(--secondary-accent);
214
+ text-decoration: none;
215
+ font-weight: 600;
216
+ transition: color .3s;
217
+ }
218
+
219
+ .container a:hover {
220
+ text-decoration: underline;
221
+ color: var(--accent-rose);
222
+ }
223
+
224
+ .container .data-box a {
225
+ position: relative;
226
+ background-image: linear-gradient(to top, var(--primary-accent), var(--primary-accent));
227
+ background-position: 0 100%;
228
+ background-repeat: no-repeat;
229
+ background-size: 0% 2px;
230
+ transition: background-size .3s, color .3s;
231
+ }
232
+
233
+ .container .data-box a:hover {
234
+ color: var(--primary-accent);
235
+ background-size: 100% 2px;
236
+ }
237
+
238
+ .container .dropdown-container {
239
+ margin-top: 20px;
240
+ }
241
+
242
+ .container .dropdown-summary {
243
+ cursor: pointer;
244
+ padding: 10px 0;
245
+ border-bottom: 1px dashed rgba(184, 194, 208, 0.4);
246
+ color: var(--text-muted);
247
+ font-size: 1.1rem;
248
+ font-weight: 400;
249
+ text-transform: uppercase;
250
+ letter-spacing: 1px;
251
+ font-family: var(--font-heading);
252
+ list-style: none;
253
+ display: flex;
254
+ align-items: center;
255
+ }
256
+
257
+ .container .dropdown-summary::-webkit-details-marker {
258
+ display: none;
259
+ }
260
+
261
+ .container .dropdown-arrow {
262
+ color: var(--primary-accent);
263
+ margin-right: 10px;
264
+ transition: transform 0.3s ease;
265
+ }
266
+
267
+ .container details[open] .dropdown-arrow {
268
+ transform: rotate(90deg);
269
+ }
270
+
271
+ .container .dropdown-content {
272
+ margin-top: 15px;
273
+ padding: 15px;
274
+ background-color: rgba(26, 29, 46, 0.6);
275
+ border-left: 2px solid var(--primary-accent);
276
+ box-shadow: 0 2px 10px rgba(236, 131, 177, 0.05);
277
+ }
278
+
279
+ .container .config-title {
280
+ color: var(--text-muted);
281
+ font-size: 1rem;
282
+ margin-bottom: 10px;
283
+ font-family: var(--font-heading);
284
+ text-transform: uppercase;
285
+ letter-spacing: 1px;
286
+ }
287
+
288
+ .container pre {
289
+ background-color: var(--bg-main);
290
+ padding: 15px;
291
+ border: 1px solid rgba(134, 197, 229, 0.4);
292
+ white-space: pre-wrap;
293
+ word-wrap: break-word;
294
+ color: var(--text-main);
295
+ border-radius: 4px;
296
+ }
297
+
298
+ .container code {
299
+ font-family: var(--font-code);
300
+ background: transparent;
301
+ padding: 0;
302
+ }
303
+ </style>
304
+ <html lang="en">
305
+ <head>
306
+ <meta charset="UTF-8">
307
+ <meta name="viewport" content="width=device-width, initial-scale=1.0">
308
+ <title>Painted Fantasy</title>
309
+ <link rel="preconnect" href="https://fonts.googleapis.com">
310
+ <link rel="preconnect" href="https://fonts.gstatic.com" crossorigin>
311
+ <link href="https://fonts.googleapis.com/css2?family=Inter:wght@300;400;600;700&family=JetBrains+Mono:wght@400;700&display=swap" rel="stylesheet">
312
+ </head>
313
+ <body>
314
+
315
+ <div class="container">
316
+ <div class="title-container">
317
+ <div class="glitchy-overlay"></div>
318
+ <div class="title-wrapper">
319
+ <h1 class="title-main">
320
+ <span class="title-prefix">PAINTED FANTASY</span>
321
+ <span class="lemonade-text">VISAGE v2</span>
322
+ </h1>
323
+ <div class="title-subtitle">
324
+ <span class="subtitle-text">Mistrall Small 3.2 Upscaled 33B</span>
325
+ </div>
326
+ </div>
327
+ </div>
328
+
329
+ ![image/png](https://cdn-uploads.huggingface.co/production/uploads/65b19c6c638328850e12d38c/9FaT1DO7_b0B_sea1LiP7.png)
330
+
331
+ <div class="section-container">
332
+ <div class="section-header">
333
+ <div class="section-indicator"></div>
334
+ <h2 class="section-title">Overview</h2>
335
+ </div>
336
+ <div class="section-content">
337
+ <p>A surprisingly difficult model to work with. Removing the repetition was coming at the expense of the unique creativity the original upscale had.</p>
338
+ <p>Decided on upscaling Painted Fantasy v2, healing it and then merging the original upscale back in.</p>
339
+ <p>The result is a smarter, uncensored, creative model that excels at character driven RP / ERP where characters are portrayed creatively and proactively.</p>
340
+ </div>
341
+ </div>
342
+
343
+ <div class="section-container">
344
+ <div class="section-header">
345
+ <div class="section-indicator"></div>
346
+ <h2 class="section-title">SillyTavern Settings</h2>
347
+ </div>
348
+ <div class="section-content">
349
+ <h3 class="subheading">Recommended Roleplay Format</h3>
350
+ <div class="data-box">
351
+ <div class="data-row">
352
+ <span class="data-arrow">></span>
353
+ <span class="data-label">Actions:</span>
354
+ <span>In plaintext</span>
355
+ </div>
356
+ <div class="data-row">
357
+ <span class="data-arrow">></span>
358
+ <span class="data-label">Dialogue:</span>
359
+ <span>"In quotes"</span>
360
+ </div>
361
+ <div class="data-row">
362
+ <span class="data-arrow">></span>
363
+ <span class="data-label">Thoughts:</span>
364
+ <span>*In asterisks*</span>
365
+ </div>
366
+ </div>
367
+ <h3 class="subheading">Recommended Samplers</h3>
368
+ <div class="data-box">
369
+ <div class="data-row">
370
+ <span class="data-arrow">></span>
371
+ <span class="data-label">Temp:</span>
372
+ <span>0.6</span>
373
+ </div>
374
+ <div class="data-row">
375
+ <span class="data-arrow">></span>
376
+ <span class="data-label">MinP:</span>
377
+ <span>0.05 - 0.1</span>
378
+ </div>
379
+ <div class="data-row">
380
+ <span class="data-arrow">></span>
381
+ <span class="data-label">TopP:</span>
382
+ <span>0.9 - 1.0</span>
383
+ </div>
384
+ <div class="data-row">
385
+ <span class="data-arrow">></span>
386
+ <span class="data-label">Dry:</span>
387
+ <span>0.8, 1.75, 4</span>
388
+ </div>
389
+ </div>
390
+ <h3 class="subheading">Instruct</h3>
391
+ <div class="data-box">
392
+ <p style="margin: 0;">Mistral v7 Tekken</p>
393
+ </div>
394
+ </div>
395
+ </div>
396
+
397
+ <div class="section-container">
398
+ <div class="section-header">
399
+ <div class="section-indicator"></div>
400
+ <h2 class="section-title">Quantizations</h2>
401
+ </div>
402
+ <div class="section-content">
403
+ <div style="margin-bottom: 20px;">
404
+ <h3 class="subheading">GGUF</h3>
405
+ <div class="data-box">
406
+ <div class="data-row">
407
+ <span class="data-arrow">></span>
408
+ <a href="https://huggingface.co/bartowski/zerofata_MS3.2-PaintedFantasy-Visage-v2-33B-GGUF">iMatrix (bartowski)</a>
409
+ </div>
410
+ </div>
411
+ </div>
412
+ <div>
413
+ <h3 class="subheading">EXL3</h3>
414
+ <div class="data-box">
415
+ <div class="data-row">
416
+ <span class="data-arrow">></span>
417
+ <a href="https://huggingface.co/zerofata/MS3.2-PaintedFantasy-Visage-v2-33B-exl3-3bpw">3bpw</a>
418
+ </div>
419
+ <div class="data-row">
420
+ <span class="data-arrow">></span>
421
+ <a href="https://huggingface.co/zerofata/MS3.2-PaintedFantasy-Visage-v2-33B-exl3-4bpw">4bpw</a>
422
+ </div>
423
+ <div class="data-row">
424
+ <span class="data-arrow">></span>
425
+ <a href="https://huggingface.co/zerofata/MS3.2-PaintedFantasy-Visage-v2-33B-exl3-5bpw">5bpw</a>
426
+ </div>
427
+ <div class="data-row">
428
+ <span class="data-arrow">></span>
429
+ <a href="https://huggingface.co/zerofata/MS3.2-PaintedFantasy-Visage-v2-33B-exl3-6bpw">6bpw</a>
430
+ </div>
431
+ </div>
432
+ </div>
433
+ </div>
434
+ </div>
435
+
436
+ <div class="section-container">
437
+ <div class="section-header">
438
+ <div class="section-indicator"></div>
439
+ <h2 class="section-title">Creation Process</h2>
440
+ </div>
441
+ <div class="section-content">
442
+ <p>Creation Process: Upscale > PT > SFT > KTO > DPO</p>
443
+ <p>Pretrained on approx 300MB of light novels, stories and FineWeb-2 corpus.</p>
444
+ <p>SFT on approx 8 million tokens, SFW / NSFW RP, stories and creative instruct data.</p>
445
+ <p>KTO on antirep data created from the SFT datasets. Rejected examples generated by MS3.2 with repetition_penalty=0.9 and OOC commands encouraging it to misgender, impersonate user etc.</p>
446
+ <p>DPO on a high quality RP / NSFW dataset that is unreleased using rejected samples created in the same method as KTO.</p>
447
+ <p>Resulting model was non repetitive, but had lost some of the spark the original upscale had. Merged the original upscale back in, making sure to not reintroduce repetition.</p>
448
+ <div class="dropdown-container">
449
+ <details>
450
+ <summary class="dropdown-summary">
451
+ <span class="dropdown-arrow">></span>
452
+ Mergekit configs
453
+ </summary>
454
+ <div class="dropdown-content">
455
+ <p>Merge configurations used during the model creation process.</p>
456
+ <div class="config-title">Initial Upscale (Passthrough)</div>
457
+ <pre><code>base_model: zerofata/MS3.2-PaintedFantasy-v2-24B
458
+ <br>
459
+ merge_method: passthrough
460
+ <br>
461
+ dtype: bfloat16
462
+ slices:
463
+ - sources:
464
+ - model: zerofata/MS3.2-PaintedFantasy-v2-24B
465
+ layer_range: [0, 29]
466
+ - sources:
467
+ - model: zerofata/MS3.2-PaintedFantasy-v2-24B
468
+ layer_range: [10, 39]</code></pre>
469
+ <div class="config-title">Final Merge (Slerp)</div>
470
+ <pre><code>models:
471
+ - model: zerofata/MS3.2-PaintedFantasy-Visage-33B
472
+ - model: ../axolotl/Visage-V2-PT-1-SFT-2-KTO-1-DPO-1/merged
473
+ merge_method: slerp
474
+ base_model: ../axolotl/Visage-V2-PT-1-SFT-2-KTO-1-DPO-1/merged
475
+ parameters:
476
+ t: [0.4, 0.2, 0, 0.2, 0.4]
477
+ dtype: bfloat16</code></pre>
478
+ </div>
479
+ </details>
480
+ </div>
481
+ <div class="dropdown-container">
482
+ <details>
483
+ <summary class="dropdown-summary">
484
+ <span class="dropdown-arrow">></span>
485
+ Axolotl configs
486
+ </summary>
487
+ <div class="dropdown-content">
488
+ <p>Not optimized for cost / performance efficiency, YMMV.</p>
489
+ <div class="config-title">Pretrain 4*H100</div>
490
+ <pre><code>&#35; ====================
491
+ &#35; MODEL CONFIGURATION
492
+ &#35; ====================
493
+ base_model: ../mergekit/pf_v2_upscale
494
+ model_type: MistralForCausalLM
495
+ tokenizer_type: AutoTokenizer
496
+ chat_template: mistral_v7_tekken
497
+ &#35; ====================
498
+ &#35; DATASET CONFIGURATION
499
+ &#35; ====================
500
+ datasets:
501
+ - path: ./data/pretrain_dataset_v5_stripped.jsonl
502
+ type: completion
503
+ <br>
504
+ dataset_prepared_path:
505
+ train_on_inputs: false &#35; Only train on assistant responses
506
+ <br>
507
+ &#35; ====================
508
+ &#35; QLORA CONFIGURATION
509
+ &#35; ====================
510
+ adapter: qlora
511
+ load_in_4bit: true
512
+ lora_r: 32
513
+ lora_alpha: 64
514
+ lora_dropout: 0.05
515
+ lora_target_linear: true
516
+ &#35; lora_modules_to_save: &#35; Uncomment only if you added NEW tokens
517
+ <br>
518
+ &#35; ====================
519
+ &#35; TRAINING PARAMETERS
520
+ &#35; ====================
521
+ num_epochs: 1
522
+ micro_batch_size: 4
523
+ gradient_accumulation_steps: 1
524
+ learning_rate: 4e-5
525
+ optimizer: paged_adamw_8bit
526
+ lr_scheduler: rex
527
+ warmup_ratio: 0.05
528
+ weight_decay: 0.01
529
+ max_grad_norm: 1.0
530
+ <br>
531
+ &#35; ====================
532
+ &#35; SEQUENCE &amp; PACKING
533
+ &#35; ====================
534
+ sequence_len: 12288
535
+ sample_packing: true
536
+ eval_sample_packing: false
537
+ pad_to_sequence_len: true
538
+ <br>
539
+ &#35; ====================
540
+ &#35; HARDWARE OPTIMIZATIONS
541
+ &#35; ====================
542
+ bf16: auto
543
+ flash_attention: true
544
+ gradient_checkpointing: offload
545
+ deepspeed: deepspeed_configs/zero1.json
546
+ <br>
547
+ plugins:
548
+ - axolotl.integrations.liger.LigerPlugin
549
+ - axolotl.integrations.cut_cross_entropy.CutCrossEntropyPlugin
550
+ cut_cross_entropy: true
551
+ liger_rope: true
552
+ liger_rms_norm: true
553
+ liger_layer_norm: true
554
+ liger_glu_activation: true
555
+ liger_cross_entropy: false &#35; Cut Cross Entropy overrides this
556
+ liger_fused_linear_cross_entropy: false &#35; Cut Cross Entropy overrides this
557
+ <br>
558
+ &#35; ====================
559
+ &#35; EVALUATION &amp; CHECKPOINTING
560
+ &#35; ====================
561
+ save_strategy: steps
562
+ save_steps: 40
563
+ save_total_limit: 5 &#35; Keep best + last few checkpoints
564
+ load_best_model_at_end: true
565
+ greater_is_better: false
566
+ <br>
567
+ &#35; ====================
568
+ &#35; LOGGING &amp; OUTPUT
569
+ &#35; ====================
570
+ output_dir: ./Visage-V2-PT-1
571
+ logging_steps: 2
572
+ save_safetensors: true
573
+ <br>
574
+ &#35; ====================
575
+ &#35; WANDB TRACKING
576
+ &#35; ====================
577
+ wandb_project: Visage-V2-PT
578
+ &#35; wandb_entity: your_entity
579
+ wandb_name: Visage-V2-PT-1</code></pre>
580
+ <div class="config-title">SFT 4*H100</div>
581
+ <pre><code>&#35; ====================
582
+ &#35; MODEL CONFIGURATION
583
+ &#35; ====================
584
+ base_model: ./Visage-V2-PT-1/merged
585
+ model_type: MistralForCausalLM
586
+ tokenizer_type: AutoTokenizer
587
+ chat_template: mistral_v7_tekken
588
+ <br>
589
+ &#35; ====================
590
+ &#35; DATASET CONFIGURATION
591
+ &#35; ====================
592
+ datasets:
593
+ - path: ./data/automated_dataset.jsonl
594
+ type: chat_template
595
+ split: train
596
+ chat_template_strategy: tokenizer
597
+ field_messages: messages
598
+ message_property_mappings:
599
+ role: role
600
+ content: content
601
+ roles:
602
+ user: ["user"]
603
+ assistant: ["assistant"]
604
+ system: ["system"]
605
+ - path: ./data/handcrafted_dataset.jsonl
606
+ type: chat_template
607
+ split: train
608
+ chat_template_strategy: tokenizer
609
+ field_messages: messages
610
+ message_property_mappings:
611
+ role: role
612
+ content: content
613
+ roles:
614
+ user: ["user"]
615
+ assistant: ["assistant"]
616
+ system: ["system"]
617
+ - path: ./data/instruct_dataset.jsonl
618
+ type: chat_template
619
+ split: train
620
+ chat_template_strategy: tokenizer
621
+ field_messages: messages
622
+ message_property_mappings:
623
+ role: role
624
+ content: content
625
+ roles:
626
+ user: ["user"]
627
+ assistant: ["assistant"]
628
+ system: ["system"]
629
+ - path: ./data/cw_dataset.jsonl
630
+ type: chat_template
631
+ split: train
632
+ chat_template_strategy: tokenizer
633
+ field_messages: messages
634
+ message_property_mappings:
635
+ role: role
636
+ content: content
637
+ roles:
638
+ user: ["user"]
639
+ assistant: ["assistant"]
640
+ system: ["system"]
641
+ - path: ./data/stories_dataset.jsonl
642
+ type: chat_template
643
+ split: train
644
+ chat_template_strategy: tokenizer
645
+ field_messages: messages
646
+ message_property_mappings:
647
+ role: role
648
+ content: content
649
+ roles:
650
+ user: ["user"]
651
+ assistant: ["assistant"]
652
+ system: ["system"]
653
+ - path: ./data/cw_claude_dataset.jsonl
654
+ type: chat_template
655
+ split: train
656
+ chat_template_strategy: tokenizer
657
+ field_messages: messages
658
+ message_property_mappings:
659
+ role: role
660
+ content: content
661
+ roles:
662
+ user: ["user"]
663
+ assistant: ["assistant"]
664
+ system: ["system"]
665
+ - path: ./data/summaries_dataset.jsonl
666
+ type: chat_template
667
+ split: train
668
+ chat_template_strategy: tokenizer
669
+ field_messages: messages
670
+ message_property_mappings:
671
+ role: role
672
+ content: content
673
+ roles:
674
+ user: ["user"]
675
+ assistant: ["assistant"]
676
+ system: ["system"]
677
+ <br>
678
+ dataset_prepared_path:
679
+ train_on_inputs: false &#35; Only train on assistant responses
680
+ <br>
681
+ &#35; ====================
682
+ &#35; QLORA CONFIGURATION
683
+ &#35; ====================
684
+ adapter: qlora
685
+ load_in_4bit: true
686
+ lora_r: 128
687
+ lora_alpha: 128
688
+ lora_dropout: 0.1
689
+ lora_target_linear: true
690
+ &#35; lora_modules_to_save: &#35; Uncomment only if you added NEW tokens
691
+ <br>
692
+ &#35; ====================
693
+ &#35; TRAINING PARAMETERS
694
+ &#35; ====================
695
+ num_epochs: 2
696
+ micro_batch_size: 2
697
+ gradient_accumulation_steps: 1
698
+ learning_rate: 1e-5
699
+ optimizer: paged_adamw_8bit
700
+ lr_scheduler: rex
701
+ warmup_ratio: 0.05
702
+ weight_decay: 0.01
703
+ max_grad_norm: 1.0
704
+ <br>
705
+ &#35; ====================
706
+ &#35; SEQUENCE &amp; PACKING
707
+ &#35; ====================
708
+ sequence_len: 8192
709
+ sample_packing: true
710
+ pad_to_sequence_len: true
711
+ <br>
712
+ &#35; ====================
713
+ &#35; HARDWARE OPTIMIZATIONS
714
+ &#35; ====================
715
+ bf16: auto
716
+ flash_attention: true
717
+ gradient_checkpointing: offload
718
+ deepspeed: deepspeed_configs/zero1.json
719
+ <br>
720
+ plugins:
721
+ - axolotl.integrations.liger.LigerPlugin
722
+ - axolotl.integrations.cut_cross_entropy.CutCrossEntropyPlugin
723
+ cut_cross_entropy: true
724
+ liger_rope: true
725
+ liger_rms_norm: true
726
+ liger_layer_norm: true
727
+ liger_glu_activation: true
728
+ liger_cross_entropy: false &#35; Cut Cross Entropy overrides this
729
+ liger_fused_linear_cross_entropy: false &#35; Cut Cross Entropy overrides this
730
+ <br>
731
+ <br>
732
+ &#35; ====================
733
+ &#35; EVALUATION &amp; CHECKPOINTING
734
+ &#35; ====================
735
+ save_strategy: steps
736
+ save_steps: 20
737
+ save_total_limit: 5 &#35; Keep best + last few checkpoints
738
+ load_best_model_at_end: true
739
+ metric_for_best_model: eval_loss
740
+ greater_is_better: false
741
+ <br>
742
+ &#35; ====================
743
+ &#35; LOGGING &amp; OUTPUT
744
+ &#35; ====================
745
+ output_dir: ./Visage-V2-PT-1-SFT-2
746
+ logging_steps: 2
747
+ save_safetensors: true
748
+ <br>
749
+ &#35; ====================
750
+ &#35; WANDB TRACKING
751
+ &#35; ====================
752
+ wandb_project: Visage-V2-SFT
753
+ &#35; wandb_entity: your_entity
754
+ wandb_name: Visage-V2-PT-1-SFT-2</code></pre>
755
+ <div class="config-title">KTO 4*H100</div>
756
+ <pre><code>&#35; ====================
757
+ &#35; MODEL CONFIGURATION
758
+ &#35; ====================
759
+ base_model: ./Visage-V2-PT-1-SFT-2/merged
760
+ model_type: MistralForCausalLM
761
+ tokenizer_type: AutoTokenizer
762
+ chat_template: mistral_v7_tekken
763
+ <br>
764
+ &#35; ====================
765
+ &#35; RL/DPO CONFIGURATION
766
+ &#35; ====================
767
+ rl: kto
768
+ rl_beta: 0.1
769
+ kto_desirable_weight: 1.25
770
+ kto_undesirable_weight: 1.0
771
+ <br>
772
+ &#35; ====================
773
+ &#35; DATASET CONFIGURATION
774
+ &#35; ====================
775
+ datasets:
776
+ - path: ./handcrafted_dataset_kto.jsonl
777
+ type: llama3.argilla
778
+ - path: ./approved_rp_dataset_kto.jsonl
779
+ type: llama3.argilla
780
+ - path: ./instruct_dataset_kto.jsonl
781
+ type: llama3.argilla
782
+ dataset_prepared_path:
783
+ train_on_inputs: false &#35; Only train on assistant responses
784
+ remove_unused_columns: False
785
+ <br>
786
+ &#35; ====================
787
+ &#35; QLORA CONFIGURATION
788
+ &#35; ====================
789
+ adapter: qlora
790
+ load_in_4bit: true
791
+ lora_r: 32
792
+ lora_alpha: 32
793
+ lora_dropout: 0.05
794
+ lora_target_linear: true
795
+ &#35; lora_modules_to_save: &#35; Uncomment only if you added NEW tokens
796
+ <br>
797
+ &#35; ====================
798
+ &#35; TRAINING PARAMETERS
799
+ &#35; ====================
800
+ num_epochs: 1
801
+ micro_batch_size: 4
802
+ gradient_accumulation_steps: 4
803
+ learning_rate: 5e-6
804
+ optimizer: adamw_8bit
805
+ lr_scheduler: cosine
806
+ warmup_steps: 15
807
+ weight_decay: 0.001
808
+ max_grad_norm: 0.01
809
+ <br>
810
+ &#35; ====================
811
+ &#35; SEQUENCE CONFIGURATION
812
+ &#35; ====================
813
+ sequence_len: 8192
814
+ pad_to_sequence_len: true
815
+ <br>
816
+ &#35; ====================
817
+ &#35; HARDWARE OPTIMIZATIONS
818
+ &#35; ====================
819
+ bf16: auto
820
+ tf32: false
821
+ flash_attention: true
822
+ gradient_checkpointing: offload
823
+ deepspeed: deepspeed_configs/zero1.json
824
+ <br>
825
+ plugins:
826
+ - axolotl.integrations.liger.LigerPlugin
827
+ - axolotl.integrations.cut_cross_entropy.CutCrossEntropyPlugin
828
+ cut_cross_entropy: true
829
+ liger_rope: true
830
+ liger_rms_norm: true
831
+ liger_layer_norm: true
832
+ liger_glu_activation: true
833
+ liger_cross_entropy: false &#35; Cut Cross Entropy overrides this
834
+ liger_fused_linear_cross_entropy: false &#35; Cut Cross Entropy overrides this
835
+ <br>
836
+ &#35; ====================
837
+ &#35; CHECKPOINTING
838
+ &#35; ====================
839
+ save_steps: 100
840
+ save_total_limit: 10
841
+ load_best_model_at_end: true
842
+ metric_for_best_model: eval_loss
843
+ greater_is_better: false
844
+ <br>
845
+ &#35; ====================
846
+ &#35; LOGGING &amp; OUTPUT
847
+ &#35; ====================
848
+ output_dir: ./Visage-V2-PT-1-SFT-2-KTO-1
849
+ logging_steps: 2
850
+ save_safetensors: true
851
+ <br>
852
+ &#35; ====================
853
+ &#35; WANDB TRACKING
854
+ &#35; ====================
855
+ wandb_project: Visage-V2-KTO
856
+ &#35; wandb_entity: your_entity
857
+ wandb_name: Visage-V2-PT-1-SFT-2-KTO-1</code></pre>
858
+ <div class="config-title">DPO 4*H100</div>
859
+ <pre><code>&#35; ====================
860
+ &#35; MODEL CONFIGURATION
861
+ &#35; ====================
862
+ base_model: ./Visage-V2-PT-1-SFT-2/merged
863
+ model_type: MistralForCausalLM
864
+ tokenizer_type: AutoTokenizer
865
+ chat_template: mistral_v7_tekken
866
+ <br>
867
+ &#35; ====================
868
+ &#35; RL/DPO CONFIGURATION
869
+ &#35; ====================
870
+ rl: dpo
871
+ rl_beta: 0.1
872
+ <br>
873
+ &#35; ====================
874
+ &#35; DATASET CONFIGURATION
875
+ &#35; ====================
876
+ datasets:
877
+ - path: ./handcrafted_dataset_mistral_rep.jsonl
878
+ type: chat_template.default
879
+ field_messages: messages
880
+ field_chosen: chosen
881
+ field_rejected: rejected
882
+ message_property_mappings:
883
+ role: role
884
+ content: content
885
+ roles:
886
+ system: ["system"]
887
+ user: ["user"]
888
+ assistant: ["assistant"]
889
+ dataset_prepared_path:
890
+ train_on_inputs: false &#35; Only train on assistant responses
891
+ <br>
892
+ &#35; ====================
893
+ &#35; QLORA CONFIGURATION
894
+ &#35; ====================
895
+ adapter: qlora
896
+ load_in_4bit: true
897
+ lora_r: 16
898
+ lora_alpha: 32
899
+ lora_dropout: 0.1
900
+ lora_target_linear: true
901
+ &#35; lora_modules_to_save: &#35; Uncomment only if you added NEW tokens
902
+ <br>
903
+ &#35; ====================
904
+ &#35; TRAINING PARAMETERS
905
+ &#35; ====================
906
+ num_epochs: 1
907
+ micro_batch_size: 2
908
+ gradient_accumulation_steps: 1
909
+ learning_rate: 2e-6
910
+ optimizer: adamw_8bit
911
+ lr_scheduler: cosine
912
+ warmup_steps: 5
913
+ weight_decay: 0.01
914
+ max_grad_norm: 1.0
915
+ <br>
916
+ &#35; ====================
917
+ &#35; SEQUENCE CONFIGURATION
918
+ &#35; ====================
919
+ sequence_len: 8192
920
+ pad_to_sequence_len: true
921
+ <br>
922
+ &#35; ====================
923
+ &#35; HARDWARE OPTIMIZATIONS
924
+ &#35; ====================
925
+ bf16: auto
926
+ tf32: false
927
+ flash_attention: true
928
+ gradient_checkpointing: offload
929
+ deepspeed: deepspeed_configs/zero1.json
930
+ <br>
931
+ plugins:
932
+ - axolotl.integrations.liger.LigerPlugin
933
+ - axolotl.integrations.cut_cross_entropy.CutCrossEntropyPlugin
934
+ cut_cross_entropy: true
935
+ liger_rope: true
936
+ liger_rms_norm: true
937
+ liger_layer_norm: true
938
+ liger_glu_activation: true
939
+ liger_cross_entropy: false &#35; Cut Cross Entropy overrides this
940
+ liger_fused_linear_cross_entropy: false &#35; Cut Cross Entropy overrides this
941
+ <br>
942
+ &#35; ====================
943
+ &#35; CHECKPOINTING
944
+ &#35; ====================
945
+ save_steps: 10
946
+ save_total_limit: 10
947
+ load_best_model_at_end: true
948
+ metric_for_best_model: eval_loss
949
+ greater_is_better: false
950
+ <br>
951
+ &#35; ====================
952
+ &#35; LOGGING &amp; OUTPUT
953
+ &#35; ====================
954
+ output_dir: ./Visage-V2-PT-1-SFT-2-DPO-1
955
+ logging_steps: 2
956
+ save_safetensors: true
957
+ <br>
958
+ &#35; ====================
959
+ &#35; WANDB TRACKING
960
+ &#35; ====================
961
+ wandb_project: Visage-V2-DPO
962
+ &#35; wandb_entity: your_entity
963
+ wandb_name: Visage-V2-PT-1-SFT-2-DPO-1</code></pre>
964
+ </div>
965
+ </details>
966
+ </div>
967
+ </div>
968
+ </div>
969
+ </div>
970
+ </body>
971
+ </html>
config.json ADDED
@@ -0,0 +1,37 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "architectures": [
3
+ "MistralForCausalLM"
4
+ ],
5
+ "attention_dropout": 0.0,
6
+ "bos_token_id": 1,
7
+ "eos_token_id": 2,
8
+ "head_dim": 128,
9
+ "hidden_act": "silu",
10
+ "hidden_size": 5120,
11
+ "initializer_range": 0.02,
12
+ "intermediate_size": 32768,
13
+ "max_position_embeddings": 131072,
14
+ "model_type": "mistral",
15
+ "num_attention_heads": 32,
16
+ "num_hidden_layers": 58,
17
+ "num_key_value_heads": 8,
18
+ "rms_norm_eps": 1e-05,
19
+ "rope_theta": 1000000000.0,
20
+ "sliding_window": null,
21
+ "tie_word_embeddings": false,
22
+ "torch_dtype": "bfloat16",
23
+ "transformers_version": "4.53.1",
24
+ "use_cache": true,
25
+ "vocab_size": 131072,
26
+ "quantization_config": {
27
+ "quant_method": "exl2",
28
+ "version": "0.3.1",
29
+ "bits": 6.0,
30
+ "head_bits": 6,
31
+ "calibration": {
32
+ "rows": 115,
33
+ "length": 2048,
34
+ "dataset": "(default)"
35
+ }
36
+ }
37
+ }
measurement.json ADDED
The diff for this file is too large to render. See raw diff
 
mergekit_config.yml ADDED
@@ -0,0 +1,8 @@
 
 
 
 
 
 
 
 
 
1
+ models:
2
+ - model: zerofata/MS3.2-PaintedFantasy-Visage-33B
3
+ - model: ../axolotl/Visage-V2-PT-1-SFT-2-KTO-1-DPO-1/merged
4
+ merge_method: slerp
5
+ base_model: ../axolotl/Visage-V2-PT-1-SFT-2-KTO-1-DPO-1/merged
6
+ parameters:
7
+ t: [0.4, 0.2, 0, 0.2, 0.4]
8
+ dtype: bfloat16
model.safetensors.index.json ADDED
@@ -0,0 +1 @@
 
 
1
+ {"metadata": {"mergekit_version": "0.1.3"}, "weight_map": {"lm_head.weight": "model-00001-of-00014.safetensors", "model.embed_tokens.weight": "model-00001-of-00014.safetensors", "model.layers.0.input_layernorm.weight": "model-00001-of-00014.safetensors", "model.layers.0.mlp.down_proj.weight": "model-00001-of-00014.safetensors", "model.layers.0.mlp.gate_proj.weight": "model-00001-of-00014.safetensors", "model.layers.0.mlp.up_proj.weight": "model-00001-of-00014.safetensors", "model.layers.0.post_attention_layernorm.weight": "model-00001-of-00014.safetensors", "model.layers.0.self_attn.k_proj.weight": "model-00001-of-00014.safetensors", "model.layers.0.self_attn.o_proj.weight": "model-00001-of-00014.safetensors", "model.layers.0.self_attn.q_proj.weight": "model-00001-of-00014.safetensors", "model.layers.0.self_attn.v_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.input_layernorm.weight": "model-00001-of-00014.safetensors", "model.layers.1.mlp.down_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.mlp.gate_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.mlp.up_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.post_attention_layernorm.weight": "model-00001-of-00014.safetensors", "model.layers.1.self_attn.k_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.self_attn.o_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.self_attn.q_proj.weight": "model-00001-of-00014.safetensors", "model.layers.1.self_attn.v_proj.weight": "model-00001-of-00014.safetensors", "model.layers.10.input_layernorm.weight": "model-00001-of-00014.safetensors", "model.layers.10.mlp.down_proj.weight": "model-00002-of-00014.safetensors", "model.layers.10.mlp.gate_proj.weight": "model-00002-of-00014.safetensors", "model.layers.10.mlp.up_proj.weight": "model-00002-of-00014.safetensors", "model.layers.10.post_attention_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.10.self_attn.k_proj.weight": "model-00002-of-00014.safetensors", "model.layers.10.self_attn.o_proj.weight": "model-00002-of-00014.safetensors", "model.layers.10.self_attn.q_proj.weight": "model-00002-of-00014.safetensors", "model.layers.10.self_attn.v_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.input_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.11.mlp.down_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.mlp.gate_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.mlp.up_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.post_attention_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.11.self_attn.k_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.self_attn.o_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.self_attn.q_proj.weight": "model-00002-of-00014.safetensors", "model.layers.11.self_attn.v_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.input_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.12.mlp.down_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.mlp.gate_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.mlp.up_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.post_attention_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.12.self_attn.k_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.self_attn.o_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.self_attn.q_proj.weight": "model-00002-of-00014.safetensors", "model.layers.12.self_attn.v_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.input_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.13.mlp.down_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.mlp.gate_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.mlp.up_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.post_attention_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.13.self_attn.k_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.self_attn.o_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.self_attn.q_proj.weight": "model-00002-of-00014.safetensors", "model.layers.13.self_attn.v_proj.weight": "model-00002-of-00014.safetensors", "model.layers.14.input_layernorm.weight": "model-00002-of-00014.safetensors", "model.layers.14.mlp.down_proj.weight": "model-00002-of-00014.safetensors", "model.layers.14.mlp.gate_proj.weight": "model-00003-of-00014.safetensors", "model.layers.14.mlp.up_proj.weight": "model-00003-of-00014.safetensors", "model.layers.14.post_attention_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.14.self_attn.k_proj.weight": "model-00003-of-00014.safetensors", "model.layers.14.self_attn.o_proj.weight": "model-00003-of-00014.safetensors", "model.layers.14.self_attn.q_proj.weight": "model-00003-of-00014.safetensors", "model.layers.14.self_attn.v_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.input_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.15.mlp.down_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.mlp.gate_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.mlp.up_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.post_attention_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.15.self_attn.k_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.self_attn.o_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.self_attn.q_proj.weight": "model-00003-of-00014.safetensors", "model.layers.15.self_attn.v_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.input_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.16.mlp.down_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.mlp.gate_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.mlp.up_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.post_attention_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.16.self_attn.k_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.self_attn.o_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.self_attn.q_proj.weight": "model-00003-of-00014.safetensors", "model.layers.16.self_attn.v_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.input_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.17.mlp.down_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.mlp.gate_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.mlp.up_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.post_attention_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.17.self_attn.k_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.self_attn.o_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.self_attn.q_proj.weight": "model-00003-of-00014.safetensors", "model.layers.17.self_attn.v_proj.weight": "model-00003-of-00014.safetensors", "model.layers.18.input_layernorm.weight": "model-00003-of-00014.safetensors", "model.layers.18.mlp.down_proj.weight": "model-00003-of-00014.safetensors", "model.layers.18.mlp.gate_proj.weight": "model-00003-of-00014.safetensors", "model.layers.18.mlp.up_proj.weight": "model-00004-of-00014.safetensors", "model.layers.18.post_attention_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.18.self_attn.k_proj.weight": "model-00004-of-00014.safetensors", "model.layers.18.self_attn.o_proj.weight": "model-00004-of-00014.safetensors", "model.layers.18.self_attn.q_proj.weight": "model-00004-of-00014.safetensors", "model.layers.18.self_attn.v_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.input_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.19.mlp.down_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.mlp.gate_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.mlp.up_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.post_attention_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.19.self_attn.k_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.self_attn.o_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.self_attn.q_proj.weight": "model-00004-of-00014.safetensors", "model.layers.19.self_attn.v_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.input_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.2.mlp.down_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.mlp.gate_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.mlp.up_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.post_attention_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.2.self_attn.k_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.self_attn.o_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.self_attn.q_proj.weight": "model-00004-of-00014.safetensors", "model.layers.2.self_attn.v_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.input_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.20.mlp.down_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.mlp.gate_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.mlp.up_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.post_attention_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.20.self_attn.k_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.self_attn.o_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.self_attn.q_proj.weight": "model-00004-of-00014.safetensors", "model.layers.20.self_attn.v_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.input_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.21.mlp.down_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.mlp.gate_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.mlp.up_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.post_attention_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.21.self_attn.k_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.self_attn.o_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.self_attn.q_proj.weight": "model-00004-of-00014.safetensors", "model.layers.21.self_attn.v_proj.weight": "model-00004-of-00014.safetensors", "model.layers.22.input_layernorm.weight": "model-00004-of-00014.safetensors", "model.layers.22.mlp.down_proj.weight": "model-00005-of-00014.safetensors", "model.layers.22.mlp.gate_proj.weight": "model-00005-of-00014.safetensors", "model.layers.22.mlp.up_proj.weight": "model-00005-of-00014.safetensors", "model.layers.22.post_attention_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.22.self_attn.k_proj.weight": "model-00005-of-00014.safetensors", "model.layers.22.self_attn.o_proj.weight": "model-00005-of-00014.safetensors", "model.layers.22.self_attn.q_proj.weight": "model-00005-of-00014.safetensors", "model.layers.22.self_attn.v_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.input_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.23.mlp.down_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.mlp.gate_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.mlp.up_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.post_attention_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.23.self_attn.k_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.self_attn.o_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.self_attn.q_proj.weight": "model-00005-of-00014.safetensors", "model.layers.23.self_attn.v_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.input_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.24.mlp.down_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.mlp.gate_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.mlp.up_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.post_attention_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.24.self_attn.k_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.self_attn.o_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.self_attn.q_proj.weight": "model-00005-of-00014.safetensors", "model.layers.24.self_attn.v_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.input_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.25.mlp.down_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.mlp.gate_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.mlp.up_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.post_attention_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.25.self_attn.k_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.self_attn.o_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.self_attn.q_proj.weight": "model-00005-of-00014.safetensors", "model.layers.25.self_attn.v_proj.weight": "model-00005-of-00014.safetensors", "model.layers.26.input_layernorm.weight": "model-00005-of-00014.safetensors", "model.layers.26.mlp.down_proj.weight": "model-00005-of-00014.safetensors", "model.layers.26.mlp.gate_proj.weight": "model-00006-of-00014.safetensors", "model.layers.26.mlp.up_proj.weight": "model-00006-of-00014.safetensors", "model.layers.26.post_attention_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.26.self_attn.k_proj.weight": "model-00006-of-00014.safetensors", "model.layers.26.self_attn.o_proj.weight": "model-00006-of-00014.safetensors", "model.layers.26.self_attn.q_proj.weight": "model-00006-of-00014.safetensors", "model.layers.26.self_attn.v_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.input_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.27.mlp.down_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.mlp.gate_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.mlp.up_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.post_attention_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.27.self_attn.k_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.self_attn.o_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.self_attn.q_proj.weight": "model-00006-of-00014.safetensors", "model.layers.27.self_attn.v_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.input_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.28.mlp.down_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.mlp.gate_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.mlp.up_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.post_attention_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.28.self_attn.k_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.self_attn.o_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.self_attn.q_proj.weight": "model-00006-of-00014.safetensors", "model.layers.28.self_attn.v_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.input_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.29.mlp.down_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.mlp.gate_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.mlp.up_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.post_attention_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.29.self_attn.k_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.self_attn.o_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.self_attn.q_proj.weight": "model-00006-of-00014.safetensors", "model.layers.29.self_attn.v_proj.weight": "model-00006-of-00014.safetensors", "model.layers.3.input_layernorm.weight": "model-00006-of-00014.safetensors", "model.layers.3.mlp.down_proj.weight": "model-00006-of-00014.safetensors", "model.layers.3.mlp.gate_proj.weight": "model-00006-of-00014.safetensors", "model.layers.3.mlp.up_proj.weight": "model-00007-of-00014.safetensors", "model.layers.3.post_attention_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.3.self_attn.k_proj.weight": "model-00007-of-00014.safetensors", "model.layers.3.self_attn.o_proj.weight": "model-00007-of-00014.safetensors", "model.layers.3.self_attn.q_proj.weight": "model-00007-of-00014.safetensors", "model.layers.3.self_attn.v_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.input_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.30.mlp.down_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.mlp.gate_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.mlp.up_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.post_attention_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.30.self_attn.k_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.self_attn.o_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.self_attn.q_proj.weight": "model-00007-of-00014.safetensors", "model.layers.30.self_attn.v_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.input_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.31.mlp.down_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.mlp.gate_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.mlp.up_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.post_attention_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.31.self_attn.k_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.self_attn.o_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.self_attn.q_proj.weight": "model-00007-of-00014.safetensors", "model.layers.31.self_attn.v_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.input_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.32.mlp.down_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.mlp.gate_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.mlp.up_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.post_attention_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.32.self_attn.k_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.self_attn.o_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.self_attn.q_proj.weight": "model-00007-of-00014.safetensors", "model.layers.32.self_attn.v_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.input_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.33.mlp.down_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.mlp.gate_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.mlp.up_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.post_attention_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.33.self_attn.k_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.self_attn.o_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.self_attn.q_proj.weight": "model-00007-of-00014.safetensors", "model.layers.33.self_attn.v_proj.weight": "model-00007-of-00014.safetensors", "model.layers.34.input_layernorm.weight": "model-00007-of-00014.safetensors", "model.layers.34.mlp.down_proj.weight": "model-00008-of-00014.safetensors", "model.layers.34.mlp.gate_proj.weight": "model-00008-of-00014.safetensors", "model.layers.34.mlp.up_proj.weight": "model-00008-of-00014.safetensors", "model.layers.34.post_attention_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.34.self_attn.k_proj.weight": "model-00008-of-00014.safetensors", "model.layers.34.self_attn.o_proj.weight": "model-00008-of-00014.safetensors", "model.layers.34.self_attn.q_proj.weight": "model-00008-of-00014.safetensors", "model.layers.34.self_attn.v_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.input_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.35.mlp.down_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.mlp.gate_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.mlp.up_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.post_attention_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.35.self_attn.k_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.self_attn.o_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.self_attn.q_proj.weight": "model-00008-of-00014.safetensors", "model.layers.35.self_attn.v_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.input_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.36.mlp.down_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.mlp.gate_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.mlp.up_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.post_attention_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.36.self_attn.k_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.self_attn.o_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.self_attn.q_proj.weight": "model-00008-of-00014.safetensors", "model.layers.36.self_attn.v_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.input_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.37.mlp.down_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.mlp.gate_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.mlp.up_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.post_attention_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.37.self_attn.k_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.self_attn.o_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.self_attn.q_proj.weight": "model-00008-of-00014.safetensors", "model.layers.37.self_attn.v_proj.weight": "model-00008-of-00014.safetensors", "model.layers.38.input_layernorm.weight": "model-00008-of-00014.safetensors", "model.layers.38.mlp.down_proj.weight": "model-00008-of-00014.safetensors", "model.layers.38.mlp.gate_proj.weight": "model-00009-of-00014.safetensors", "model.layers.38.mlp.up_proj.weight": "model-00009-of-00014.safetensors", "model.layers.38.post_attention_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.38.self_attn.k_proj.weight": "model-00009-of-00014.safetensors", "model.layers.38.self_attn.o_proj.weight": "model-00009-of-00014.safetensors", "model.layers.38.self_attn.q_proj.weight": "model-00009-of-00014.safetensors", "model.layers.38.self_attn.v_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.input_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.39.mlp.down_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.mlp.gate_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.mlp.up_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.post_attention_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.39.self_attn.k_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.self_attn.o_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.self_attn.q_proj.weight": "model-00009-of-00014.safetensors", "model.layers.39.self_attn.v_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.input_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.4.mlp.down_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.mlp.gate_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.mlp.up_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.post_attention_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.4.self_attn.k_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.self_attn.o_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.self_attn.q_proj.weight": "model-00009-of-00014.safetensors", "model.layers.4.self_attn.v_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.input_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.40.mlp.down_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.mlp.gate_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.mlp.up_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.post_attention_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.40.self_attn.k_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.self_attn.o_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.self_attn.q_proj.weight": "model-00009-of-00014.safetensors", "model.layers.40.self_attn.v_proj.weight": "model-00009-of-00014.safetensors", "model.layers.41.input_layernorm.weight": "model-00009-of-00014.safetensors", "model.layers.41.mlp.down_proj.weight": "model-00009-of-00014.safetensors", "model.layers.41.mlp.gate_proj.weight": "model-00009-of-00014.safetensors", "model.layers.41.mlp.up_proj.weight": "model-00010-of-00014.safetensors", "model.layers.41.post_attention_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.41.self_attn.k_proj.weight": "model-00010-of-00014.safetensors", "model.layers.41.self_attn.o_proj.weight": "model-00010-of-00014.safetensors", "model.layers.41.self_attn.q_proj.weight": "model-00010-of-00014.safetensors", "model.layers.41.self_attn.v_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.input_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.42.mlp.down_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.mlp.gate_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.mlp.up_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.post_attention_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.42.self_attn.k_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.self_attn.o_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.self_attn.q_proj.weight": "model-00010-of-00014.safetensors", "model.layers.42.self_attn.v_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.input_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.43.mlp.down_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.mlp.gate_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.mlp.up_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.post_attention_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.43.self_attn.k_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.self_attn.o_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.self_attn.q_proj.weight": "model-00010-of-00014.safetensors", "model.layers.43.self_attn.v_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.input_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.44.mlp.down_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.mlp.gate_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.mlp.up_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.post_attention_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.44.self_attn.k_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.self_attn.o_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.self_attn.q_proj.weight": "model-00010-of-00014.safetensors", "model.layers.44.self_attn.v_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.input_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.45.mlp.down_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.mlp.gate_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.mlp.up_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.post_attention_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.45.self_attn.k_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.self_attn.o_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.self_attn.q_proj.weight": "model-00010-of-00014.safetensors", "model.layers.45.self_attn.v_proj.weight": "model-00010-of-00014.safetensors", "model.layers.46.input_layernorm.weight": "model-00010-of-00014.safetensors", "model.layers.46.mlp.down_proj.weight": "model-00011-of-00014.safetensors", "model.layers.46.mlp.gate_proj.weight": "model-00011-of-00014.safetensors", "model.layers.46.mlp.up_proj.weight": "model-00011-of-00014.safetensors", "model.layers.46.post_attention_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.46.self_attn.k_proj.weight": "model-00011-of-00014.safetensors", "model.layers.46.self_attn.o_proj.weight": "model-00011-of-00014.safetensors", "model.layers.46.self_attn.q_proj.weight": "model-00011-of-00014.safetensors", "model.layers.46.self_attn.v_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.input_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.47.mlp.down_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.mlp.gate_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.mlp.up_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.post_attention_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.47.self_attn.k_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.self_attn.o_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.self_attn.q_proj.weight": "model-00011-of-00014.safetensors", "model.layers.47.self_attn.v_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.input_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.48.mlp.down_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.mlp.gate_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.mlp.up_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.post_attention_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.48.self_attn.k_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.self_attn.o_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.self_attn.q_proj.weight": "model-00011-of-00014.safetensors", "model.layers.48.self_attn.v_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.input_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.49.mlp.down_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.mlp.gate_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.mlp.up_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.post_attention_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.49.self_attn.k_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.self_attn.o_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.self_attn.q_proj.weight": "model-00011-of-00014.safetensors", "model.layers.49.self_attn.v_proj.weight": "model-00011-of-00014.safetensors", "model.layers.5.input_layernorm.weight": "model-00011-of-00014.safetensors", "model.layers.5.mlp.down_proj.weight": "model-00011-of-00014.safetensors", "model.layers.5.mlp.gate_proj.weight": "model-00012-of-00014.safetensors", "model.layers.5.mlp.up_proj.weight": "model-00012-of-00014.safetensors", "model.layers.5.post_attention_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.5.self_attn.k_proj.weight": "model-00012-of-00014.safetensors", "model.layers.5.self_attn.o_proj.weight": "model-00012-of-00014.safetensors", "model.layers.5.self_attn.q_proj.weight": "model-00012-of-00014.safetensors", "model.layers.5.self_attn.v_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.input_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.50.mlp.down_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.mlp.gate_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.mlp.up_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.post_attention_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.50.self_attn.k_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.self_attn.o_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.self_attn.q_proj.weight": "model-00012-of-00014.safetensors", "model.layers.50.self_attn.v_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.input_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.51.mlp.down_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.mlp.gate_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.mlp.up_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.post_attention_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.51.self_attn.k_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.self_attn.o_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.self_attn.q_proj.weight": "model-00012-of-00014.safetensors", "model.layers.51.self_attn.v_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.input_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.52.mlp.down_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.mlp.gate_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.mlp.up_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.post_attention_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.52.self_attn.k_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.self_attn.o_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.self_attn.q_proj.weight": "model-00012-of-00014.safetensors", "model.layers.52.self_attn.v_proj.weight": "model-00012-of-00014.safetensors", "model.layers.53.input_layernorm.weight": "model-00012-of-00014.safetensors", "model.layers.53.mlp.down_proj.weight": "model-00012-of-00014.safetensors", "model.layers.53.mlp.gate_proj.weight": "model-00012-of-00014.safetensors", "model.layers.53.mlp.up_proj.weight": "model-00013-of-00014.safetensors", "model.layers.53.post_attention_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.53.self_attn.k_proj.weight": "model-00013-of-00014.safetensors", "model.layers.53.self_attn.o_proj.weight": "model-00013-of-00014.safetensors", "model.layers.53.self_attn.q_proj.weight": "model-00013-of-00014.safetensors", "model.layers.53.self_attn.v_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.input_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.54.mlp.down_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.mlp.gate_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.mlp.up_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.post_attention_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.54.self_attn.k_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.self_attn.o_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.self_attn.q_proj.weight": "model-00013-of-00014.safetensors", "model.layers.54.self_attn.v_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.input_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.55.mlp.down_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.mlp.gate_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.mlp.up_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.post_attention_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.55.self_attn.k_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.self_attn.o_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.self_attn.q_proj.weight": "model-00013-of-00014.safetensors", "model.layers.55.self_attn.v_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.input_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.56.mlp.down_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.mlp.gate_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.mlp.up_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.post_attention_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.56.self_attn.k_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.self_attn.o_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.self_attn.q_proj.weight": "model-00013-of-00014.safetensors", "model.layers.56.self_attn.v_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.input_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.57.mlp.down_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.mlp.gate_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.mlp.up_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.post_attention_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.57.self_attn.k_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.self_attn.o_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.self_attn.q_proj.weight": "model-00013-of-00014.safetensors", "model.layers.57.self_attn.v_proj.weight": "model-00013-of-00014.safetensors", "model.layers.6.input_layernorm.weight": "model-00013-of-00014.safetensors", "model.layers.6.mlp.down_proj.weight": "model-00014-of-00014.safetensors", "model.layers.6.mlp.gate_proj.weight": "model-00014-of-00014.safetensors", "model.layers.6.mlp.up_proj.weight": "model-00014-of-00014.safetensors", "model.layers.6.post_attention_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.6.self_attn.k_proj.weight": "model-00014-of-00014.safetensors", "model.layers.6.self_attn.o_proj.weight": "model-00014-of-00014.safetensors", "model.layers.6.self_attn.q_proj.weight": "model-00014-of-00014.safetensors", "model.layers.6.self_attn.v_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.input_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.7.mlp.down_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.mlp.gate_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.mlp.up_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.post_attention_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.7.self_attn.k_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.self_attn.o_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.self_attn.q_proj.weight": "model-00014-of-00014.safetensors", "model.layers.7.self_attn.v_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.input_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.8.mlp.down_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.mlp.gate_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.mlp.up_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.post_attention_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.8.self_attn.k_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.self_attn.o_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.self_attn.q_proj.weight": "model-00014-of-00014.safetensors", "model.layers.8.self_attn.v_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.input_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.9.mlp.down_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.mlp.gate_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.mlp.up_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.post_attention_layernorm.weight": "model-00014-of-00014.safetensors", "model.layers.9.self_attn.k_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.self_attn.o_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.self_attn.q_proj.weight": "model-00014-of-00014.safetensors", "model.layers.9.self_attn.v_proj.weight": "model-00014-of-00014.safetensors", "model.norm.weight": "model-00014-of-00014.safetensors"}}
output-00001-of-00004.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5dd0cc62aea13f26c16184cbf8ddd75fd177abcda409bb59b6a5f06262559a23
3
+ size 8468517088
output-00002-of-00004.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:776e4d84f0fb5fc41578c9373c17cccaa87a1771e05abbf3a9c90e87204aeeb4
3
+ size 8565770074
output-00003-of-00004.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5ad1dc5b7b7dbc01f5948dc38dbb2fd7a72b6c1d0235f1fc978b0cb741ab05cd
3
+ size 8486756066
output-00004-of-00004.safetensors ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:cc8346d47a762febf2bf123fc714e69ac9078429b22103949b3b3627dc1660d2
3
+ size 528482400
special_tokens_map.json ADDED
@@ -0,0 +1,1032 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "additional_special_tokens": [
3
+ "<unk>",
4
+ "<s>",
5
+ "</s>",
6
+ "[INST]",
7
+ "[/INST]",
8
+ "[AVAILABLE_TOOLS]",
9
+ "[/AVAILABLE_TOOLS]",
10
+ "[TOOL_RESULTS]",
11
+ "[/TOOL_RESULTS]",
12
+ "[TOOL_CALLS]",
13
+ "[IMG]",
14
+ "<pad>",
15
+ "[IMG_BREAK]",
16
+ "[IMG_END]",
17
+ "[PREFIX]",
18
+ "[MIDDLE]",
19
+ "[SUFFIX]",
20
+ "[SYSTEM_PROMPT]",
21
+ "[/SYSTEM_PROMPT]",
22
+ "[TOOL_CONTENT]",
23
+ "<SPECIAL_20>",
24
+ "<SPECIAL_21>",
25
+ "<SPECIAL_22>",
26
+ "<SPECIAL_23>",
27
+ "<SPECIAL_24>",
28
+ "<SPECIAL_25>",
29
+ "<SPECIAL_26>",
30
+ "<SPECIAL_27>",
31
+ "<SPECIAL_28>",
32
+ "<SPECIAL_29>",
33
+ "<SPECIAL_30>",
34
+ "<SPECIAL_31>",
35
+ "<SPECIAL_32>",
36
+ "<SPECIAL_33>",
37
+ "<SPECIAL_34>",
38
+ "<SPECIAL_35>",
39
+ "<SPECIAL_36>",
40
+ "<SPECIAL_37>",
41
+ "<SPECIAL_38>",
42
+ "<SPECIAL_39>",
43
+ "<SPECIAL_40>",
44
+ "<SPECIAL_41>",
45
+ "<SPECIAL_42>",
46
+ "<SPECIAL_43>",
47
+ "<SPECIAL_44>",
48
+ "<SPECIAL_45>",
49
+ "<SPECIAL_46>",
50
+ "<SPECIAL_47>",
51
+ "<SPECIAL_48>",
52
+ "<SPECIAL_49>",
53
+ "<SPECIAL_50>",
54
+ "<SPECIAL_51>",
55
+ "<SPECIAL_52>",
56
+ "<SPECIAL_53>",
57
+ "<SPECIAL_54>",
58
+ "<SPECIAL_55>",
59
+ "<SPECIAL_56>",
60
+ "<SPECIAL_57>",
61
+ "<SPECIAL_58>",
62
+ "<SPECIAL_59>",
63
+ "<SPECIAL_60>",
64
+ "<SPECIAL_61>",
65
+ "<SPECIAL_62>",
66
+ "<SPECIAL_63>",
67
+ "<SPECIAL_64>",
68
+ "<SPECIAL_65>",
69
+ "<SPECIAL_66>",
70
+ "<SPECIAL_67>",
71
+ "<SPECIAL_68>",
72
+ "<SPECIAL_69>",
73
+ "<SPECIAL_70>",
74
+ "<SPECIAL_71>",
75
+ "<SPECIAL_72>",
76
+ "<SPECIAL_73>",
77
+ "<SPECIAL_74>",
78
+ "<SPECIAL_75>",
79
+ "<SPECIAL_76>",
80
+ "<SPECIAL_77>",
81
+ "<SPECIAL_78>",
82
+ "<SPECIAL_79>",
83
+ "<SPECIAL_80>",
84
+ "<SPECIAL_81>",
85
+ "<SPECIAL_82>",
86
+ "<SPECIAL_83>",
87
+ "<SPECIAL_84>",
88
+ "<SPECIAL_85>",
89
+ "<SPECIAL_86>",
90
+ "<SPECIAL_87>",
91
+ "<SPECIAL_88>",
92
+ "<SPECIAL_89>",
93
+ "<SPECIAL_90>",
94
+ "<SPECIAL_91>",
95
+ "<SPECIAL_92>",
96
+ "<SPECIAL_93>",
97
+ "<SPECIAL_94>",
98
+ "<SPECIAL_95>",
99
+ "<SPECIAL_96>",
100
+ "<SPECIAL_97>",
101
+ "<SPECIAL_98>",
102
+ "<SPECIAL_99>",
103
+ "<SPECIAL_100>",
104
+ "<SPECIAL_101>",
105
+ "<SPECIAL_102>",
106
+ "<SPECIAL_103>",
107
+ "<SPECIAL_104>",
108
+ "<SPECIAL_105>",
109
+ "<SPECIAL_106>",
110
+ "<SPECIAL_107>",
111
+ "<SPECIAL_108>",
112
+ "<SPECIAL_109>",
113
+ "<SPECIAL_110>",
114
+ "<SPECIAL_111>",
115
+ "<SPECIAL_112>",
116
+ "<SPECIAL_113>",
117
+ "<SPECIAL_114>",
118
+ "<SPECIAL_115>",
119
+ "<SPECIAL_116>",
120
+ "<SPECIAL_117>",
121
+ "<SPECIAL_118>",
122
+ "<SPECIAL_119>",
123
+ "<SPECIAL_120>",
124
+ "<SPECIAL_121>",
125
+ "<SPECIAL_122>",
126
+ "<SPECIAL_123>",
127
+ "<SPECIAL_124>",
128
+ "<SPECIAL_125>",
129
+ "<SPECIAL_126>",
130
+ "<SPECIAL_127>",
131
+ "<SPECIAL_128>",
132
+ "<SPECIAL_129>",
133
+ "<SPECIAL_130>",
134
+ "<SPECIAL_131>",
135
+ "<SPECIAL_132>",
136
+ "<SPECIAL_133>",
137
+ "<SPECIAL_134>",
138
+ "<SPECIAL_135>",
139
+ "<SPECIAL_136>",
140
+ "<SPECIAL_137>",
141
+ "<SPECIAL_138>",
142
+ "<SPECIAL_139>",
143
+ "<SPECIAL_140>",
144
+ "<SPECIAL_141>",
145
+ "<SPECIAL_142>",
146
+ "<SPECIAL_143>",
147
+ "<SPECIAL_144>",
148
+ "<SPECIAL_145>",
149
+ "<SPECIAL_146>",
150
+ "<SPECIAL_147>",
151
+ "<SPECIAL_148>",
152
+ "<SPECIAL_149>",
153
+ "<SPECIAL_150>",
154
+ "<SPECIAL_151>",
155
+ "<SPECIAL_152>",
156
+ "<SPECIAL_153>",
157
+ "<SPECIAL_154>",
158
+ "<SPECIAL_155>",
159
+ "<SPECIAL_156>",
160
+ "<SPECIAL_157>",
161
+ "<SPECIAL_158>",
162
+ "<SPECIAL_159>",
163
+ "<SPECIAL_160>",
164
+ "<SPECIAL_161>",
165
+ "<SPECIAL_162>",
166
+ "<SPECIAL_163>",
167
+ "<SPECIAL_164>",
168
+ "<SPECIAL_165>",
169
+ "<SPECIAL_166>",
170
+ "<SPECIAL_167>",
171
+ "<SPECIAL_168>",
172
+ "<SPECIAL_169>",
173
+ "<SPECIAL_170>",
174
+ "<SPECIAL_171>",
175
+ "<SPECIAL_172>",
176
+ "<SPECIAL_173>",
177
+ "<SPECIAL_174>",
178
+ "<SPECIAL_175>",
179
+ "<SPECIAL_176>",
180
+ "<SPECIAL_177>",
181
+ "<SPECIAL_178>",
182
+ "<SPECIAL_179>",
183
+ "<SPECIAL_180>",
184
+ "<SPECIAL_181>",
185
+ "<SPECIAL_182>",
186
+ "<SPECIAL_183>",
187
+ "<SPECIAL_184>",
188
+ "<SPECIAL_185>",
189
+ "<SPECIAL_186>",
190
+ "<SPECIAL_187>",
191
+ "<SPECIAL_188>",
192
+ "<SPECIAL_189>",
193
+ "<SPECIAL_190>",
194
+ "<SPECIAL_191>",
195
+ "<SPECIAL_192>",
196
+ "<SPECIAL_193>",
197
+ "<SPECIAL_194>",
198
+ "<SPECIAL_195>",
199
+ "<SPECIAL_196>",
200
+ "<SPECIAL_197>",
201
+ "<SPECIAL_198>",
202
+ "<SPECIAL_199>",
203
+ "<SPECIAL_200>",
204
+ "<SPECIAL_201>",
205
+ "<SPECIAL_202>",
206
+ "<SPECIAL_203>",
207
+ "<SPECIAL_204>",
208
+ "<SPECIAL_205>",
209
+ "<SPECIAL_206>",
210
+ "<SPECIAL_207>",
211
+ "<SPECIAL_208>",
212
+ "<SPECIAL_209>",
213
+ "<SPECIAL_210>",
214
+ "<SPECIAL_211>",
215
+ "<SPECIAL_212>",
216
+ "<SPECIAL_213>",
217
+ "<SPECIAL_214>",
218
+ "<SPECIAL_215>",
219
+ "<SPECIAL_216>",
220
+ "<SPECIAL_217>",
221
+ "<SPECIAL_218>",
222
+ "<SPECIAL_219>",
223
+ "<SPECIAL_220>",
224
+ "<SPECIAL_221>",
225
+ "<SPECIAL_222>",
226
+ "<SPECIAL_223>",
227
+ "<SPECIAL_224>",
228
+ "<SPECIAL_225>",
229
+ "<SPECIAL_226>",
230
+ "<SPECIAL_227>",
231
+ "<SPECIAL_228>",
232
+ "<SPECIAL_229>",
233
+ "<SPECIAL_230>",
234
+ "<SPECIAL_231>",
235
+ "<SPECIAL_232>",
236
+ "<SPECIAL_233>",
237
+ "<SPECIAL_234>",
238
+ "<SPECIAL_235>",
239
+ "<SPECIAL_236>",
240
+ "<SPECIAL_237>",
241
+ "<SPECIAL_238>",
242
+ "<SPECIAL_239>",
243
+ "<SPECIAL_240>",
244
+ "<SPECIAL_241>",
245
+ "<SPECIAL_242>",
246
+ "<SPECIAL_243>",
247
+ "<SPECIAL_244>",
248
+ "<SPECIAL_245>",
249
+ "<SPECIAL_246>",
250
+ "<SPECIAL_247>",
251
+ "<SPECIAL_248>",
252
+ "<SPECIAL_249>",
253
+ "<SPECIAL_250>",
254
+ "<SPECIAL_251>",
255
+ "<SPECIAL_252>",
256
+ "<SPECIAL_253>",
257
+ "<SPECIAL_254>",
258
+ "<SPECIAL_255>",
259
+ "<SPECIAL_256>",
260
+ "<SPECIAL_257>",
261
+ "<SPECIAL_258>",
262
+ "<SPECIAL_259>",
263
+ "<SPECIAL_260>",
264
+ "<SPECIAL_261>",
265
+ "<SPECIAL_262>",
266
+ "<SPECIAL_263>",
267
+ "<SPECIAL_264>",
268
+ "<SPECIAL_265>",
269
+ "<SPECIAL_266>",
270
+ "<SPECIAL_267>",
271
+ "<SPECIAL_268>",
272
+ "<SPECIAL_269>",
273
+ "<SPECIAL_270>",
274
+ "<SPECIAL_271>",
275
+ "<SPECIAL_272>",
276
+ "<SPECIAL_273>",
277
+ "<SPECIAL_274>",
278
+ "<SPECIAL_275>",
279
+ "<SPECIAL_276>",
280
+ "<SPECIAL_277>",
281
+ "<SPECIAL_278>",
282
+ "<SPECIAL_279>",
283
+ "<SPECIAL_280>",
284
+ "<SPECIAL_281>",
285
+ "<SPECIAL_282>",
286
+ "<SPECIAL_283>",
287
+ "<SPECIAL_284>",
288
+ "<SPECIAL_285>",
289
+ "<SPECIAL_286>",
290
+ "<SPECIAL_287>",
291
+ "<SPECIAL_288>",
292
+ "<SPECIAL_289>",
293
+ "<SPECIAL_290>",
294
+ "<SPECIAL_291>",
295
+ "<SPECIAL_292>",
296
+ "<SPECIAL_293>",
297
+ "<SPECIAL_294>",
298
+ "<SPECIAL_295>",
299
+ "<SPECIAL_296>",
300
+ "<SPECIAL_297>",
301
+ "<SPECIAL_298>",
302
+ "<SPECIAL_299>",
303
+ "<SPECIAL_300>",
304
+ "<SPECIAL_301>",
305
+ "<SPECIAL_302>",
306
+ "<SPECIAL_303>",
307
+ "<SPECIAL_304>",
308
+ "<SPECIAL_305>",
309
+ "<SPECIAL_306>",
310
+ "<SPECIAL_307>",
311
+ "<SPECIAL_308>",
312
+ "<SPECIAL_309>",
313
+ "<SPECIAL_310>",
314
+ "<SPECIAL_311>",
315
+ "<SPECIAL_312>",
316
+ "<SPECIAL_313>",
317
+ "<SPECIAL_314>",
318
+ "<SPECIAL_315>",
319
+ "<SPECIAL_316>",
320
+ "<SPECIAL_317>",
321
+ "<SPECIAL_318>",
322
+ "<SPECIAL_319>",
323
+ "<SPECIAL_320>",
324
+ "<SPECIAL_321>",
325
+ "<SPECIAL_322>",
326
+ "<SPECIAL_323>",
327
+ "<SPECIAL_324>",
328
+ "<SPECIAL_325>",
329
+ "<SPECIAL_326>",
330
+ "<SPECIAL_327>",
331
+ "<SPECIAL_328>",
332
+ "<SPECIAL_329>",
333
+ "<SPECIAL_330>",
334
+ "<SPECIAL_331>",
335
+ "<SPECIAL_332>",
336
+ "<SPECIAL_333>",
337
+ "<SPECIAL_334>",
338
+ "<SPECIAL_335>",
339
+ "<SPECIAL_336>",
340
+ "<SPECIAL_337>",
341
+ "<SPECIAL_338>",
342
+ "<SPECIAL_339>",
343
+ "<SPECIAL_340>",
344
+ "<SPECIAL_341>",
345
+ "<SPECIAL_342>",
346
+ "<SPECIAL_343>",
347
+ "<SPECIAL_344>",
348
+ "<SPECIAL_345>",
349
+ "<SPECIAL_346>",
350
+ "<SPECIAL_347>",
351
+ "<SPECIAL_348>",
352
+ "<SPECIAL_349>",
353
+ "<SPECIAL_350>",
354
+ "<SPECIAL_351>",
355
+ "<SPECIAL_352>",
356
+ "<SPECIAL_353>",
357
+ "<SPECIAL_354>",
358
+ "<SPECIAL_355>",
359
+ "<SPECIAL_356>",
360
+ "<SPECIAL_357>",
361
+ "<SPECIAL_358>",
362
+ "<SPECIAL_359>",
363
+ "<SPECIAL_360>",
364
+ "<SPECIAL_361>",
365
+ "<SPECIAL_362>",
366
+ "<SPECIAL_363>",
367
+ "<SPECIAL_364>",
368
+ "<SPECIAL_365>",
369
+ "<SPECIAL_366>",
370
+ "<SPECIAL_367>",
371
+ "<SPECIAL_368>",
372
+ "<SPECIAL_369>",
373
+ "<SPECIAL_370>",
374
+ "<SPECIAL_371>",
375
+ "<SPECIAL_372>",
376
+ "<SPECIAL_373>",
377
+ "<SPECIAL_374>",
378
+ "<SPECIAL_375>",
379
+ "<SPECIAL_376>",
380
+ "<SPECIAL_377>",
381
+ "<SPECIAL_378>",
382
+ "<SPECIAL_379>",
383
+ "<SPECIAL_380>",
384
+ "<SPECIAL_381>",
385
+ "<SPECIAL_382>",
386
+ "<SPECIAL_383>",
387
+ "<SPECIAL_384>",
388
+ "<SPECIAL_385>",
389
+ "<SPECIAL_386>",
390
+ "<SPECIAL_387>",
391
+ "<SPECIAL_388>",
392
+ "<SPECIAL_389>",
393
+ "<SPECIAL_390>",
394
+ "<SPECIAL_391>",
395
+ "<SPECIAL_392>",
396
+ "<SPECIAL_393>",
397
+ "<SPECIAL_394>",
398
+ "<SPECIAL_395>",
399
+ "<SPECIAL_396>",
400
+ "<SPECIAL_397>",
401
+ "<SPECIAL_398>",
402
+ "<SPECIAL_399>",
403
+ "<SPECIAL_400>",
404
+ "<SPECIAL_401>",
405
+ "<SPECIAL_402>",
406
+ "<SPECIAL_403>",
407
+ "<SPECIAL_404>",
408
+ "<SPECIAL_405>",
409
+ "<SPECIAL_406>",
410
+ "<SPECIAL_407>",
411
+ "<SPECIAL_408>",
412
+ "<SPECIAL_409>",
413
+ "<SPECIAL_410>",
414
+ "<SPECIAL_411>",
415
+ "<SPECIAL_412>",
416
+ "<SPECIAL_413>",
417
+ "<SPECIAL_414>",
418
+ "<SPECIAL_415>",
419
+ "<SPECIAL_416>",
420
+ "<SPECIAL_417>",
421
+ "<SPECIAL_418>",
422
+ "<SPECIAL_419>",
423
+ "<SPECIAL_420>",
424
+ "<SPECIAL_421>",
425
+ "<SPECIAL_422>",
426
+ "<SPECIAL_423>",
427
+ "<SPECIAL_424>",
428
+ "<SPECIAL_425>",
429
+ "<SPECIAL_426>",
430
+ "<SPECIAL_427>",
431
+ "<SPECIAL_428>",
432
+ "<SPECIAL_429>",
433
+ "<SPECIAL_430>",
434
+ "<SPECIAL_431>",
435
+ "<SPECIAL_432>",
436
+ "<SPECIAL_433>",
437
+ "<SPECIAL_434>",
438
+ "<SPECIAL_435>",
439
+ "<SPECIAL_436>",
440
+ "<SPECIAL_437>",
441
+ "<SPECIAL_438>",
442
+ "<SPECIAL_439>",
443
+ "<SPECIAL_440>",
444
+ "<SPECIAL_441>",
445
+ "<SPECIAL_442>",
446
+ "<SPECIAL_443>",
447
+ "<SPECIAL_444>",
448
+ "<SPECIAL_445>",
449
+ "<SPECIAL_446>",
450
+ "<SPECIAL_447>",
451
+ "<SPECIAL_448>",
452
+ "<SPECIAL_449>",
453
+ "<SPECIAL_450>",
454
+ "<SPECIAL_451>",
455
+ "<SPECIAL_452>",
456
+ "<SPECIAL_453>",
457
+ "<SPECIAL_454>",
458
+ "<SPECIAL_455>",
459
+ "<SPECIAL_456>",
460
+ "<SPECIAL_457>",
461
+ "<SPECIAL_458>",
462
+ "<SPECIAL_459>",
463
+ "<SPECIAL_460>",
464
+ "<SPECIAL_461>",
465
+ "<SPECIAL_462>",
466
+ "<SPECIAL_463>",
467
+ "<SPECIAL_464>",
468
+ "<SPECIAL_465>",
469
+ "<SPECIAL_466>",
470
+ "<SPECIAL_467>",
471
+ "<SPECIAL_468>",
472
+ "<SPECIAL_469>",
473
+ "<SPECIAL_470>",
474
+ "<SPECIAL_471>",
475
+ "<SPECIAL_472>",
476
+ "<SPECIAL_473>",
477
+ "<SPECIAL_474>",
478
+ "<SPECIAL_475>",
479
+ "<SPECIAL_476>",
480
+ "<SPECIAL_477>",
481
+ "<SPECIAL_478>",
482
+ "<SPECIAL_479>",
483
+ "<SPECIAL_480>",
484
+ "<SPECIAL_481>",
485
+ "<SPECIAL_482>",
486
+ "<SPECIAL_483>",
487
+ "<SPECIAL_484>",
488
+ "<SPECIAL_485>",
489
+ "<SPECIAL_486>",
490
+ "<SPECIAL_487>",
491
+ "<SPECIAL_488>",
492
+ "<SPECIAL_489>",
493
+ "<SPECIAL_490>",
494
+ "<SPECIAL_491>",
495
+ "<SPECIAL_492>",
496
+ "<SPECIAL_493>",
497
+ "<SPECIAL_494>",
498
+ "<SPECIAL_495>",
499
+ "<SPECIAL_496>",
500
+ "<SPECIAL_497>",
501
+ "<SPECIAL_498>",
502
+ "<SPECIAL_499>",
503
+ "<SPECIAL_500>",
504
+ "<SPECIAL_501>",
505
+ "<SPECIAL_502>",
506
+ "<SPECIAL_503>",
507
+ "<SPECIAL_504>",
508
+ "<SPECIAL_505>",
509
+ "<SPECIAL_506>",
510
+ "<SPECIAL_507>",
511
+ "<SPECIAL_508>",
512
+ "<SPECIAL_509>",
513
+ "<SPECIAL_510>",
514
+ "<SPECIAL_511>",
515
+ "<SPECIAL_512>",
516
+ "<SPECIAL_513>",
517
+ "<SPECIAL_514>",
518
+ "<SPECIAL_515>",
519
+ "<SPECIAL_516>",
520
+ "<SPECIAL_517>",
521
+ "<SPECIAL_518>",
522
+ "<SPECIAL_519>",
523
+ "<SPECIAL_520>",
524
+ "<SPECIAL_521>",
525
+ "<SPECIAL_522>",
526
+ "<SPECIAL_523>",
527
+ "<SPECIAL_524>",
528
+ "<SPECIAL_525>",
529
+ "<SPECIAL_526>",
530
+ "<SPECIAL_527>",
531
+ "<SPECIAL_528>",
532
+ "<SPECIAL_529>",
533
+ "<SPECIAL_530>",
534
+ "<SPECIAL_531>",
535
+ "<SPECIAL_532>",
536
+ "<SPECIAL_533>",
537
+ "<SPECIAL_534>",
538
+ "<SPECIAL_535>",
539
+ "<SPECIAL_536>",
540
+ "<SPECIAL_537>",
541
+ "<SPECIAL_538>",
542
+ "<SPECIAL_539>",
543
+ "<SPECIAL_540>",
544
+ "<SPECIAL_541>",
545
+ "<SPECIAL_542>",
546
+ "<SPECIAL_543>",
547
+ "<SPECIAL_544>",
548
+ "<SPECIAL_545>",
549
+ "<SPECIAL_546>",
550
+ "<SPECIAL_547>",
551
+ "<SPECIAL_548>",
552
+ "<SPECIAL_549>",
553
+ "<SPECIAL_550>",
554
+ "<SPECIAL_551>",
555
+ "<SPECIAL_552>",
556
+ "<SPECIAL_553>",
557
+ "<SPECIAL_554>",
558
+ "<SPECIAL_555>",
559
+ "<SPECIAL_556>",
560
+ "<SPECIAL_557>",
561
+ "<SPECIAL_558>",
562
+ "<SPECIAL_559>",
563
+ "<SPECIAL_560>",
564
+ "<SPECIAL_561>",
565
+ "<SPECIAL_562>",
566
+ "<SPECIAL_563>",
567
+ "<SPECIAL_564>",
568
+ "<SPECIAL_565>",
569
+ "<SPECIAL_566>",
570
+ "<SPECIAL_567>",
571
+ "<SPECIAL_568>",
572
+ "<SPECIAL_569>",
573
+ "<SPECIAL_570>",
574
+ "<SPECIAL_571>",
575
+ "<SPECIAL_572>",
576
+ "<SPECIAL_573>",
577
+ "<SPECIAL_574>",
578
+ "<SPECIAL_575>",
579
+ "<SPECIAL_576>",
580
+ "<SPECIAL_577>",
581
+ "<SPECIAL_578>",
582
+ "<SPECIAL_579>",
583
+ "<SPECIAL_580>",
584
+ "<SPECIAL_581>",
585
+ "<SPECIAL_582>",
586
+ "<SPECIAL_583>",
587
+ "<SPECIAL_584>",
588
+ "<SPECIAL_585>",
589
+ "<SPECIAL_586>",
590
+ "<SPECIAL_587>",
591
+ "<SPECIAL_588>",
592
+ "<SPECIAL_589>",
593
+ "<SPECIAL_590>",
594
+ "<SPECIAL_591>",
595
+ "<SPECIAL_592>",
596
+ "<SPECIAL_593>",
597
+ "<SPECIAL_594>",
598
+ "<SPECIAL_595>",
599
+ "<SPECIAL_596>",
600
+ "<SPECIAL_597>",
601
+ "<SPECIAL_598>",
602
+ "<SPECIAL_599>",
603
+ "<SPECIAL_600>",
604
+ "<SPECIAL_601>",
605
+ "<SPECIAL_602>",
606
+ "<SPECIAL_603>",
607
+ "<SPECIAL_604>",
608
+ "<SPECIAL_605>",
609
+ "<SPECIAL_606>",
610
+ "<SPECIAL_607>",
611
+ "<SPECIAL_608>",
612
+ "<SPECIAL_609>",
613
+ "<SPECIAL_610>",
614
+ "<SPECIAL_611>",
615
+ "<SPECIAL_612>",
616
+ "<SPECIAL_613>",
617
+ "<SPECIAL_614>",
618
+ "<SPECIAL_615>",
619
+ "<SPECIAL_616>",
620
+ "<SPECIAL_617>",
621
+ "<SPECIAL_618>",
622
+ "<SPECIAL_619>",
623
+ "<SPECIAL_620>",
624
+ "<SPECIAL_621>",
625
+ "<SPECIAL_622>",
626
+ "<SPECIAL_623>",
627
+ "<SPECIAL_624>",
628
+ "<SPECIAL_625>",
629
+ "<SPECIAL_626>",
630
+ "<SPECIAL_627>",
631
+ "<SPECIAL_628>",
632
+ "<SPECIAL_629>",
633
+ "<SPECIAL_630>",
634
+ "<SPECIAL_631>",
635
+ "<SPECIAL_632>",
636
+ "<SPECIAL_633>",
637
+ "<SPECIAL_634>",
638
+ "<SPECIAL_635>",
639
+ "<SPECIAL_636>",
640
+ "<SPECIAL_637>",
641
+ "<SPECIAL_638>",
642
+ "<SPECIAL_639>",
643
+ "<SPECIAL_640>",
644
+ "<SPECIAL_641>",
645
+ "<SPECIAL_642>",
646
+ "<SPECIAL_643>",
647
+ "<SPECIAL_644>",
648
+ "<SPECIAL_645>",
649
+ "<SPECIAL_646>",
650
+ "<SPECIAL_647>",
651
+ "<SPECIAL_648>",
652
+ "<SPECIAL_649>",
653
+ "<SPECIAL_650>",
654
+ "<SPECIAL_651>",
655
+ "<SPECIAL_652>",
656
+ "<SPECIAL_653>",
657
+ "<SPECIAL_654>",
658
+ "<SPECIAL_655>",
659
+ "<SPECIAL_656>",
660
+ "<SPECIAL_657>",
661
+ "<SPECIAL_658>",
662
+ "<SPECIAL_659>",
663
+ "<SPECIAL_660>",
664
+ "<SPECIAL_661>",
665
+ "<SPECIAL_662>",
666
+ "<SPECIAL_663>",
667
+ "<SPECIAL_664>",
668
+ "<SPECIAL_665>",
669
+ "<SPECIAL_666>",
670
+ "<SPECIAL_667>",
671
+ "<SPECIAL_668>",
672
+ "<SPECIAL_669>",
673
+ "<SPECIAL_670>",
674
+ "<SPECIAL_671>",
675
+ "<SPECIAL_672>",
676
+ "<SPECIAL_673>",
677
+ "<SPECIAL_674>",
678
+ "<SPECIAL_675>",
679
+ "<SPECIAL_676>",
680
+ "<SPECIAL_677>",
681
+ "<SPECIAL_678>",
682
+ "<SPECIAL_679>",
683
+ "<SPECIAL_680>",
684
+ "<SPECIAL_681>",
685
+ "<SPECIAL_682>",
686
+ "<SPECIAL_683>",
687
+ "<SPECIAL_684>",
688
+ "<SPECIAL_685>",
689
+ "<SPECIAL_686>",
690
+ "<SPECIAL_687>",
691
+ "<SPECIAL_688>",
692
+ "<SPECIAL_689>",
693
+ "<SPECIAL_690>",
694
+ "<SPECIAL_691>",
695
+ "<SPECIAL_692>",
696
+ "<SPECIAL_693>",
697
+ "<SPECIAL_694>",
698
+ "<SPECIAL_695>",
699
+ "<SPECIAL_696>",
700
+ "<SPECIAL_697>",
701
+ "<SPECIAL_698>",
702
+ "<SPECIAL_699>",
703
+ "<SPECIAL_700>",
704
+ "<SPECIAL_701>",
705
+ "<SPECIAL_702>",
706
+ "<SPECIAL_703>",
707
+ "<SPECIAL_704>",
708
+ "<SPECIAL_705>",
709
+ "<SPECIAL_706>",
710
+ "<SPECIAL_707>",
711
+ "<SPECIAL_708>",
712
+ "<SPECIAL_709>",
713
+ "<SPECIAL_710>",
714
+ "<SPECIAL_711>",
715
+ "<SPECIAL_712>",
716
+ "<SPECIAL_713>",
717
+ "<SPECIAL_714>",
718
+ "<SPECIAL_715>",
719
+ "<SPECIAL_716>",
720
+ "<SPECIAL_717>",
721
+ "<SPECIAL_718>",
722
+ "<SPECIAL_719>",
723
+ "<SPECIAL_720>",
724
+ "<SPECIAL_721>",
725
+ "<SPECIAL_722>",
726
+ "<SPECIAL_723>",
727
+ "<SPECIAL_724>",
728
+ "<SPECIAL_725>",
729
+ "<SPECIAL_726>",
730
+ "<SPECIAL_727>",
731
+ "<SPECIAL_728>",
732
+ "<SPECIAL_729>",
733
+ "<SPECIAL_730>",
734
+ "<SPECIAL_731>",
735
+ "<SPECIAL_732>",
736
+ "<SPECIAL_733>",
737
+ "<SPECIAL_734>",
738
+ "<SPECIAL_735>",
739
+ "<SPECIAL_736>",
740
+ "<SPECIAL_737>",
741
+ "<SPECIAL_738>",
742
+ "<SPECIAL_739>",
743
+ "<SPECIAL_740>",
744
+ "<SPECIAL_741>",
745
+ "<SPECIAL_742>",
746
+ "<SPECIAL_743>",
747
+ "<SPECIAL_744>",
748
+ "<SPECIAL_745>",
749
+ "<SPECIAL_746>",
750
+ "<SPECIAL_747>",
751
+ "<SPECIAL_748>",
752
+ "<SPECIAL_749>",
753
+ "<SPECIAL_750>",
754
+ "<SPECIAL_751>",
755
+ "<SPECIAL_752>",
756
+ "<SPECIAL_753>",
757
+ "<SPECIAL_754>",
758
+ "<SPECIAL_755>",
759
+ "<SPECIAL_756>",
760
+ "<SPECIAL_757>",
761
+ "<SPECIAL_758>",
762
+ "<SPECIAL_759>",
763
+ "<SPECIAL_760>",
764
+ "<SPECIAL_761>",
765
+ "<SPECIAL_762>",
766
+ "<SPECIAL_763>",
767
+ "<SPECIAL_764>",
768
+ "<SPECIAL_765>",
769
+ "<SPECIAL_766>",
770
+ "<SPECIAL_767>",
771
+ "<SPECIAL_768>",
772
+ "<SPECIAL_769>",
773
+ "<SPECIAL_770>",
774
+ "<SPECIAL_771>",
775
+ "<SPECIAL_772>",
776
+ "<SPECIAL_773>",
777
+ "<SPECIAL_774>",
778
+ "<SPECIAL_775>",
779
+ "<SPECIAL_776>",
780
+ "<SPECIAL_777>",
781
+ "<SPECIAL_778>",
782
+ "<SPECIAL_779>",
783
+ "<SPECIAL_780>",
784
+ "<SPECIAL_781>",
785
+ "<SPECIAL_782>",
786
+ "<SPECIAL_783>",
787
+ "<SPECIAL_784>",
788
+ "<SPECIAL_785>",
789
+ "<SPECIAL_786>",
790
+ "<SPECIAL_787>",
791
+ "<SPECIAL_788>",
792
+ "<SPECIAL_789>",
793
+ "<SPECIAL_790>",
794
+ "<SPECIAL_791>",
795
+ "<SPECIAL_792>",
796
+ "<SPECIAL_793>",
797
+ "<SPECIAL_794>",
798
+ "<SPECIAL_795>",
799
+ "<SPECIAL_796>",
800
+ "<SPECIAL_797>",
801
+ "<SPECIAL_798>",
802
+ "<SPECIAL_799>",
803
+ "<SPECIAL_800>",
804
+ "<SPECIAL_801>",
805
+ "<SPECIAL_802>",
806
+ "<SPECIAL_803>",
807
+ "<SPECIAL_804>",
808
+ "<SPECIAL_805>",
809
+ "<SPECIAL_806>",
810
+ "<SPECIAL_807>",
811
+ "<SPECIAL_808>",
812
+ "<SPECIAL_809>",
813
+ "<SPECIAL_810>",
814
+ "<SPECIAL_811>",
815
+ "<SPECIAL_812>",
816
+ "<SPECIAL_813>",
817
+ "<SPECIAL_814>",
818
+ "<SPECIAL_815>",
819
+ "<SPECIAL_816>",
820
+ "<SPECIAL_817>",
821
+ "<SPECIAL_818>",
822
+ "<SPECIAL_819>",
823
+ "<SPECIAL_820>",
824
+ "<SPECIAL_821>",
825
+ "<SPECIAL_822>",
826
+ "<SPECIAL_823>",
827
+ "<SPECIAL_824>",
828
+ "<SPECIAL_825>",
829
+ "<SPECIAL_826>",
830
+ "<SPECIAL_827>",
831
+ "<SPECIAL_828>",
832
+ "<SPECIAL_829>",
833
+ "<SPECIAL_830>",
834
+ "<SPECIAL_831>",
835
+ "<SPECIAL_832>",
836
+ "<SPECIAL_833>",
837
+ "<SPECIAL_834>",
838
+ "<SPECIAL_835>",
839
+ "<SPECIAL_836>",
840
+ "<SPECIAL_837>",
841
+ "<SPECIAL_838>",
842
+ "<SPECIAL_839>",
843
+ "<SPECIAL_840>",
844
+ "<SPECIAL_841>",
845
+ "<SPECIAL_842>",
846
+ "<SPECIAL_843>",
847
+ "<SPECIAL_844>",
848
+ "<SPECIAL_845>",
849
+ "<SPECIAL_846>",
850
+ "<SPECIAL_847>",
851
+ "<SPECIAL_848>",
852
+ "<SPECIAL_849>",
853
+ "<SPECIAL_850>",
854
+ "<SPECIAL_851>",
855
+ "<SPECIAL_852>",
856
+ "<SPECIAL_853>",
857
+ "<SPECIAL_854>",
858
+ "<SPECIAL_855>",
859
+ "<SPECIAL_856>",
860
+ "<SPECIAL_857>",
861
+ "<SPECIAL_858>",
862
+ "<SPECIAL_859>",
863
+ "<SPECIAL_860>",
864
+ "<SPECIAL_861>",
865
+ "<SPECIAL_862>",
866
+ "<SPECIAL_863>",
867
+ "<SPECIAL_864>",
868
+ "<SPECIAL_865>",
869
+ "<SPECIAL_866>",
870
+ "<SPECIAL_867>",
871
+ "<SPECIAL_868>",
872
+ "<SPECIAL_869>",
873
+ "<SPECIAL_870>",
874
+ "<SPECIAL_871>",
875
+ "<SPECIAL_872>",
876
+ "<SPECIAL_873>",
877
+ "<SPECIAL_874>",
878
+ "<SPECIAL_875>",
879
+ "<SPECIAL_876>",
880
+ "<SPECIAL_877>",
881
+ "<SPECIAL_878>",
882
+ "<SPECIAL_879>",
883
+ "<SPECIAL_880>",
884
+ "<SPECIAL_881>",
885
+ "<SPECIAL_882>",
886
+ "<SPECIAL_883>",
887
+ "<SPECIAL_884>",
888
+ "<SPECIAL_885>",
889
+ "<SPECIAL_886>",
890
+ "<SPECIAL_887>",
891
+ "<SPECIAL_888>",
892
+ "<SPECIAL_889>",
893
+ "<SPECIAL_890>",
894
+ "<SPECIAL_891>",
895
+ "<SPECIAL_892>",
896
+ "<SPECIAL_893>",
897
+ "<SPECIAL_894>",
898
+ "<SPECIAL_895>",
899
+ "<SPECIAL_896>",
900
+ "<SPECIAL_897>",
901
+ "<SPECIAL_898>",
902
+ "<SPECIAL_899>",
903
+ "<SPECIAL_900>",
904
+ "<SPECIAL_901>",
905
+ "<SPECIAL_902>",
906
+ "<SPECIAL_903>",
907
+ "<SPECIAL_904>",
908
+ "<SPECIAL_905>",
909
+ "<SPECIAL_906>",
910
+ "<SPECIAL_907>",
911
+ "<SPECIAL_908>",
912
+ "<SPECIAL_909>",
913
+ "<SPECIAL_910>",
914
+ "<SPECIAL_911>",
915
+ "<SPECIAL_912>",
916
+ "<SPECIAL_913>",
917
+ "<SPECIAL_914>",
918
+ "<SPECIAL_915>",
919
+ "<SPECIAL_916>",
920
+ "<SPECIAL_917>",
921
+ "<SPECIAL_918>",
922
+ "<SPECIAL_919>",
923
+ "<SPECIAL_920>",
924
+ "<SPECIAL_921>",
925
+ "<SPECIAL_922>",
926
+ "<SPECIAL_923>",
927
+ "<SPECIAL_924>",
928
+ "<SPECIAL_925>",
929
+ "<SPECIAL_926>",
930
+ "<SPECIAL_927>",
931
+ "<SPECIAL_928>",
932
+ "<SPECIAL_929>",
933
+ "<SPECIAL_930>",
934
+ "<SPECIAL_931>",
935
+ "<SPECIAL_932>",
936
+ "<SPECIAL_933>",
937
+ "<SPECIAL_934>",
938
+ "<SPECIAL_935>",
939
+ "<SPECIAL_936>",
940
+ "<SPECIAL_937>",
941
+ "<SPECIAL_938>",
942
+ "<SPECIAL_939>",
943
+ "<SPECIAL_940>",
944
+ "<SPECIAL_941>",
945
+ "<SPECIAL_942>",
946
+ "<SPECIAL_943>",
947
+ "<SPECIAL_944>",
948
+ "<SPECIAL_945>",
949
+ "<SPECIAL_946>",
950
+ "<SPECIAL_947>",
951
+ "<SPECIAL_948>",
952
+ "<SPECIAL_949>",
953
+ "<SPECIAL_950>",
954
+ "<SPECIAL_951>",
955
+ "<SPECIAL_952>",
956
+ "<SPECIAL_953>",
957
+ "<SPECIAL_954>",
958
+ "<SPECIAL_955>",
959
+ "<SPECIAL_956>",
960
+ "<SPECIAL_957>",
961
+ "<SPECIAL_958>",
962
+ "<SPECIAL_959>",
963
+ "<SPECIAL_960>",
964
+ "<SPECIAL_961>",
965
+ "<SPECIAL_962>",
966
+ "<SPECIAL_963>",
967
+ "<SPECIAL_964>",
968
+ "<SPECIAL_965>",
969
+ "<SPECIAL_966>",
970
+ "<SPECIAL_967>",
971
+ "<SPECIAL_968>",
972
+ "<SPECIAL_969>",
973
+ "<SPECIAL_970>",
974
+ "<SPECIAL_971>",
975
+ "<SPECIAL_972>",
976
+ "<SPECIAL_973>",
977
+ "<SPECIAL_974>",
978
+ "<SPECIAL_975>",
979
+ "<SPECIAL_976>",
980
+ "<SPECIAL_977>",
981
+ "<SPECIAL_978>",
982
+ "<SPECIAL_979>",
983
+ "<SPECIAL_980>",
984
+ "<SPECIAL_981>",
985
+ "<SPECIAL_982>",
986
+ "<SPECIAL_983>",
987
+ "<SPECIAL_984>",
988
+ "<SPECIAL_985>",
989
+ "<SPECIAL_986>",
990
+ "<SPECIAL_987>",
991
+ "<SPECIAL_988>",
992
+ "<SPECIAL_989>",
993
+ "<SPECIAL_990>",
994
+ "<SPECIAL_991>",
995
+ "<SPECIAL_992>",
996
+ "<SPECIAL_993>",
997
+ "<SPECIAL_994>",
998
+ "<SPECIAL_995>",
999
+ "<SPECIAL_996>",
1000
+ "<SPECIAL_997>",
1001
+ "<SPECIAL_998>",
1002
+ "<SPECIAL_999>"
1003
+ ],
1004
+ "bos_token": {
1005
+ "content": "<s>",
1006
+ "lstrip": false,
1007
+ "normalized": false,
1008
+ "rstrip": false,
1009
+ "single_word": false
1010
+ },
1011
+ "eos_token": {
1012
+ "content": "</s>",
1013
+ "lstrip": false,
1014
+ "normalized": false,
1015
+ "rstrip": false,
1016
+ "single_word": false
1017
+ },
1018
+ "pad_token": {
1019
+ "content": "<pad>",
1020
+ "lstrip": false,
1021
+ "normalized": false,
1022
+ "rstrip": false,
1023
+ "single_word": false
1024
+ },
1025
+ "unk_token": {
1026
+ "content": "<unk>",
1027
+ "lstrip": false,
1028
+ "normalized": false,
1029
+ "rstrip": false,
1030
+ "single_word": false
1031
+ }
1032
+ }
tokenizer.json ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:b76085f9923309d873994d444989f7eb6ec074b06f25b58f1e8d7b7741070949
3
+ size 17078037
tokenizer_config.json ADDED
The diff for this file is too large to render. See raw diff