@misc{indiciae579b86197d50, title = {CLAPSep: Leveraging Contrastive Pre-trained Model for Multi-Modal Query-Conditioned Target Sound Extraction}, author = {Hao Ma and Zhiyuan Peng and Xu Li and Mingjie Shao and Xixin Wu and Ju Liu}, year = {2025}, url = {https://arxiv.org/abs/2402.17455}, note = {Source identifier: 2402.17455} }