@misc{indiciae3cdc24925864, title = {MMAudioSep: Taming Video-to-Audio Generative Model Towards Video/Text-Queried Sound Separation}, author = {Akira Takahashi and Shusuke Takahashi and Yuki Mitsufuji}, year = {2026}, url = {https://arxiv.org/abs/2510.09065}, note = {Source identifier: 2510.09065} }