@misc{indiciae12fbba7f390c, title = {Speech Rhythm-Based Speaker Embeddings Extraction from Phonemes and Phoneme Duration for Multi-Speaker Speech Synthesis}, author = {Kenichi Fujita and Atsushi Ando and Yusuke Ijima}, year = {2024}, doi = {10.1587/transinf.2023edp7039}, url = {https://arxiv.org/abs/2402.07085}, note = {Source identifier: 2402.07085} }