Views
No views yet
pip install TensorFlowTTS1import numpy as np
2import soundfile as sf
3import yaml
4
5import tensorflow as tf
6
7from tensorflow_tts.inference import AutoProcessor
8from tensorflow_tts.inference import TFAutoModel
9
10processor = AutoProcessor.from_pretrained("tensorspeech/tts-tacotron2-thorsten-ger")
11tacotron2 = TFAutoModel.from_pretrained("tensorspeech/tts-tacotron2-thorsten-ger")
12
13text = "Möchtest du das meiner Frau erklären? Nein? Ich auch nicht."
14
15input_ids = processor.text_to_sequence(text)
16
17decoder_output, mel_outputs, stop_token_prediction, alignment_history = tacotron2.inference(
18 input_ids=tf.expand_dims(tf.convert_to_tensor(input_ids, dtype=tf.int32), 0),
19 input_lengths=tf.convert_to_tensor([len(input_ids)], tf.int32),
20 speaker_ids=tf.convert_to_tensor([0], dtype=tf.int32),
21)
22@article{DBLP:journals/corr/abs-1712-05884,
author = {Jonathan Shen and
Ruoming Pang and
Ron J. Weiss and
Mike Schuster and
Navdeep Jaitly and
Zongheng Yang and
Zhifeng Chen and
Yu Zhang and
Yuxuan Wang and
R. J. Skerry{-}Ryan and
Rif A. Saurous and
Yannis Agiomyrgiannakis and
Yonghui Wu},
title = {Natural {TTS} Synthesis by Conditioning WaveNet on Mel Spectrogram
Predictions},
journal = {CoRR},
volume = {abs/1712.05884},
year = {2017},
url = {http://arxiv.org/abs/1712.05884},
archivePrefix = {arXiv},
eprint = {1712.05884},
timestamp = {Thu, 28 Nov 2019 08:59:52 +0100},
biburl = {https://dblp.org/rec/journals/corr/abs-1712-05884.bib},
bibsource = {dblp computer science bibliography, https://dblp.org}
}@misc{TFTTS,
author = {Minh Nguyen, Alejandro Miguel Velasquez, Erogol, Kuan Chen, Dawid Kobus, Takuya Ebata,
Trinh Le and Yunchao He},
title = {TensorflowTTS},
year = {2020},
publisher = {GitHub},
journal = {GitHub repository},
howpublished = {\\url{https://github.com/TensorSpeech/TensorFlowTTS}},
}