@misc{indiciaea1c7b5eeb3ea, title = {VoiceWeaver: Staged Learning of Structured Controls for Expressive Speech and Sound-Event Generation}, author = {Xiaosu Su and Yun Cao and Yiping Ni and Xiaowei Yi}, year = {2026}, url = {https://arxiv.org/abs/2610.04500}, note = {Source identifier: 2610.04500} }