jonatasgrosman
commited on
Commit
•
02a9eb3
1
Parent(s):
9688a91
add LM
Browse files- alphabet.json +1 -0
- config.json +0 -1
- language_model/attrs.json +1 -0
- language_model/lm.binary +3 -0
- language_model/unigrams.txt +3 -0
- preprocessor_config.json +2 -1
- vocab.json +1 -1
alphabet.json
ADDED
@@ -0,0 +1 @@
|
|
|
|
|
1 |
+
{"labels": ["", "<s>", "</s>", "⁇", " ", "'", "-", "a", "b", "c", "d", "e", "f", "g", "h", "i", "j", "k", "l", "m", "n", "o", "p", "q", "r", "s", "t", "u", "v", "w", "x", "y", "z", "ä", "í", "ó", "ö", "ü"], "is_bpe": false}
|
config.json
CHANGED
@@ -48,7 +48,6 @@
|
|
48 |
"feat_proj_dropout": 0.05,
|
49 |
"feat_quantizer_dropout": 0.0,
|
50 |
"final_dropout": 0.0,
|
51 |
-
"gradient_checkpointing": true,
|
52 |
"hidden_act": "gelu",
|
53 |
"hidden_dropout": 0.05,
|
54 |
"hidden_size": 1024,
|
|
|
48 |
"feat_proj_dropout": 0.05,
|
49 |
"feat_quantizer_dropout": 0.0,
|
50 |
"final_dropout": 0.0,
|
|
|
51 |
"hidden_act": "gelu",
|
52 |
"hidden_dropout": 0.05,
|
53 |
"hidden_size": 1024,
|
language_model/attrs.json
ADDED
@@ -0,0 +1 @@
|
|
|
|
|
1 |
+
{"alpha": 0.5, "beta": 1.5, "unk_score_offset": -10.0, "score_boundary": true}
|
language_model/lm.binary
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1b406a87efac267f6404fe991869186701f40ef2ea2e6304a97eaa1cb928e7eb
|
3 |
+
size 1393421477
|
language_model/unigrams.txt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:ade813c30a291ddb95c434014f882c013f0c216cd4e6dedc4bee425c223c58d5
|
3 |
+
size 28506089
|
preprocessor_config.json
CHANGED
@@ -5,5 +5,6 @@
|
|
5 |
"padding_side": "right",
|
6 |
"padding_value": 0.0,
|
7 |
"return_attention_mask": true,
|
8 |
-
"sampling_rate": 16000
|
|
|
9 |
}
|
|
|
5 |
"padding_side": "right",
|
6 |
"padding_value": 0.0,
|
7 |
"return_attention_mask": true,
|
8 |
+
"sampling_rate": 16000,
|
9 |
+
"processor_class": "Wav2Vec2ProcessorWithLM"
|
10 |
}
|
vocab.json
CHANGED
@@ -1 +1 @@
|
|
1 |
-
{"<pad>": 0, "<s>": 1, "</s>": 2, "<unk>": 3, "|": 4, "'": 5, "-": 6, "
|
|
|
1 |
+
{"<pad>": 0, "<s>": 1, "</s>": 2, "<unk>": 3, "|": 4, "'": 5, "-": 6, "a": 7, "b": 8, "c": 9, "d": 10, "e": 11, "f": 12, "g": 13, "h": 14, "i": 15, "j": 16, "k": 17, "l": 18, "m": 19, "n": 20, "o": 21, "p": 22, "q": 23, "r": 24, "s": 25, "t": 26, "u": 27, "v": 28, "w": 29, "x": 30, "y": 31, "z": 32, "ä": 33, "í": 34, "ó": 35, "ö": 36, "ü": 37}
|