@misc{indiciaeb2b48bed2a11, title = {CATR: Combinatorial-Dependence Audio-Queried Transformer for Audio-Visual Video Segmentation}, author = {Kexin Li and Zongxin Yang and Lei Chen and Yi Yang and Jun Xiao}, year = {2023}, url = {https://arxiv.org/abs/2309.09709}, note = {Source identifier: 2309.09709} }