@misc{indiciaec24ddf7b6667, title = {MF-Speech: Achieving Fine-Grained and Compositional Control in Speech Generation via Factor Disentanglement}, author = {Xinyue Yu and Youqing Fang and Pingyu Wu and Guoyang Ye and Wenbo Zhou and Weiming Zhang and Song Xiao}, year = {2025}, url = {https://arxiv.org/abs/2511.12074}, note = {Source identifier: 2511.12074} }