wav2vec2-60-urdu / vocab.json
kingabzpro's picture
add tokenizer
00918c5
raw
history blame
588 Bytes
{"،": 1, "؟": 2, "ء": 3, "آ": 4, "ؤ": 5, "ئ": 6, "ا": 7, "ب": 8, "ت": 9, "ث": 10, "ج": 11, "ح": 12, "خ": 13, "د": 14, "ذ": 15, "ر": 16, "ز": 17, "س": 18, "ش": 19, "ص": 20, "ض": 21, "ط": 22, "ظ": 23, "ع": 24, "غ": 25, "ف": 26, "ق": 27, "ل": 28, "م": 29, "ن": 30, "و": 31, "ى": 32, "ي": 33, "َ": 34, "ُ": 35, "ِ": 36, "ّ": 37, "ٔ": 38, "ٰ": 39, "ٹ": 40, "پ": 41, "چ": 42, "ڈ": 43, "ڑ": 44, "ژ": 45, "ک": 46, "گ": 47, "ں": 48, "ھ": 49, "ہ": 50, "ۂ": 51, "ی": 52, "ے": 53, "۔": 54, "|": 0, "<unk>": 55, "<pad>": 56, "<s>": 57, "</s>": 58}