Views
No views yet
pip install nanospeechpython -m nanospeech.generate --text "The quick brown fox jumps over the lazy dog."--voice parameter to select the voice used for speech:celeste — Sampleluna — Samplenash — Sampleorion — Samplerhea — Samplepython -m nanospeech.generate --help for a full list of options to customize the voice.1@article{chen-etal-2024-f5tts,
2 title = {F5-TTS: A Fairytaler that Fakes Fluent and Faithful Speech with Flow Matching},
3 author = {Yushen Chen and Zhikang Niu and Ziyang Ma and Keqi Deng and Chunhui Wang and Jian Zhao and Kai Yu and Xie Chen},
4 year = {2024},
5 url = {https://api.semanticscholar.org/CorpusID:273228169}
6}1@inproceedings{Eskimez2024E2TE,
2 title = {E2 TTS: Embarrassingly Easy Fully Non-Autoregressive Zero-Shot TTS},
3 author = {Sefik Emre Eskimez and Xiaofei Wang and Manthan Thakker and Canrun Li and Chung-Hsien Tsai and Zhen Xiao and Hemin Yang and Zirun Zhu and Min Tang and Xu Tan and Yanqing Liu and Sheng Zhao and Naoyuki Kanda},
4 year = {2024},
5 url = {https://api.semanticscholar.org/CorpusID:270738197}
6}1@article{Le2023VoiceboxTM,
2 title = {Voicebox: Text-Guided Multilingual Universal Speech Generation at Scale},
3 author = {Matt Le and Apoorv Vyas and Bowen Shi and Brian Karrer and Leda Sari and Rashel Moritz and Mary Williamson and Vimal Manohar and Yossi Adi and Jay Mahadeokar and Wei-Ning Hsu},
4 year = {2023},
5 url = {https://api.semanticscholar.org/CorpusID:259275061}
6}1@article{tong2023generalized,
2 title = {Improving and Generalizing Flow-Based Generative Models with Minibatch Optimal Transport},
3 author = {Alexander Tong and Joshua Fan and Ricky T. Q. Chen and Jesse Bettencourt and David Duvenaud},
4 year = {2023}
5 url = {https://api.semanticscholar.org/CorpusID:259847293}
6}1@article{peebles2022scalable,
2 title = {Scalable Diffusion Models with Transformers},
3 author = {Peebles, William and Xie, Saining},
4 year = {2022},
5 url = {https://api.semanticscholar.org/CorpusID:254854389}
6}1@article{lipman2022flow,
2 title = {Flow Matching for Generative Modeling},
3 author = {Yaron Lipman and Ricky T. Q. Chen and Heli Ben-Hamu and Maximilian Nickel and Matt Le},
4 year = {2022},
5 url = {https://api.semanticscholar.org/CorpusID:252734897}
6}1@article{koizumi2023librittsr,
2 title = {LibriTTS-R: A Restored Multi-Speaker Text-to-Speech Corpus},
3 author = {Yuma Koizumi and Heiga Zen and Shigeki Karita and Yifan Ding and Kohei Yatabe and Nobuyuki Morioka and Michiel Bacchiani and Yu Zhang and Wei Han and Ankur Bapna},
4 year = {2023},
5 url = {https://api.semanticscholar.org/CorpusID:258967444}
6}