@misc{indiciae66e18bcf74ee, title = {DurFlex-EVC: Duration-Flexible Emotional Voice Conversion Leveraging Discrete Representations without Text Alignment}, author = {Hyung-Seok Oh and Sang-Hoon Lee and Deok-Hyeon Cho and Seong-Whan Lee}, year = {2025}, doi = {10.1109/taffc.2025.3530920}, url = {https://arxiv.org/abs/2401.08095}, note = {Source identifier: 2401.08095} }