File size: 1,974 Bytes
1182a51
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
102
103
104
105
106
107
108
109
110
111
112
113
114
115
116
117
118
119
120
121
122
123
124
125
126
127
128
129
130
131
132
133
134
135
136
137
138
139
140
141
142
143
144
145
{
  "data": {
    "text_cleaners": [
      "cjke_cleaners2"
    ],
    "sampling_rate": 22050,
    "filter_length": 1024,
    "hop_length": 256,
    "win_length": 1024,
    "n_mel_channels": 80,
    "add_blank": true,
    "cleaned_text": true,
    "n_speakers": 10
  },
  "model": {
    "inter_channels": 192,
    "hidden_channels": 192,
    "filter_channels": 768,
    "n_heads": 2,
    "n_layers": 6,
    "n_layers_trans_flow": 3,
    "kernel_size": 3,
    "p_dropout": 0.1,
    "resblock": "1",
    "resblock_kernel_sizes": [
      3,
      7,
      11
    ],
    "resblock_dilation_sizes": [
      [
        1,
        3,
        5
      ],
      [
        1,
        3,
        5
      ],
      [
        1,
        3,
        5
      ]
    ],
    "upsample_rates": [
      8,
      8,
      2,
      2
    ],
    "upsample_initial_channel": 512,
    "upsample_kernel_sizes": [
      16,
      16,
      4,
      4
    ],
    "n_layers_q": 3,
    "use_spectral_norm": false,
    "gin_channels": 256
  },
  "symbols": [
    "_",
    ",",
    ".",
    "!",
    "?",
    "-",
    "~",
    "\u2026",
    "N",
    "Q",
    "a",
    "b",
    "d",
    "e",
    "f",
    "g",
    "h",
    "i",
    "j",
    "k",
    "l",
    "m",
    "n",
    "o",
    "p",
    "s",
    "t",
    "u",
    "v",
    "w",
    "x",
    "y",
    "z",
    "\u0251",
    "\u00e6",
    "\u0283",
    "\u0291",
    "\u00e7",
    "\u026f",
    "\u026a",
    "\u0254",
    "\u025b",
    "\u0279",
    "\u00f0",
    "\u0259",
    "\u026b",
    "\u0265",
    "\u0278",
    "\u028a",
    "\u027e",
    "\u0292",
    "\u03b8",
    "\u03b2",
    "\u014b",
    "\u0266",
    "\u207c",
    "\u02b0",
    "`",
    "^",
    "#",
    "*",
    "=",
    "\u02c8",
    "\u02cc",
    "\u2192",
    "\u2193",
    "\u2191",
    " "
  ],
  "speakers": {
    "default": 1,
    "whispering": 2,
    "shouting": 3,
    "excited": 4,
    "cheerful": 5,
    "terrified": 6,
    "angry": 7,
    "sad": 8,
    "friendly": 9
  }
}