@misc{indiciae8c78dc6aff59, title = {Bridging Audio and Vision: Zero-Shot Audiovisual Segmentation by Connecting Pretrained Models}, author = {Seung-jae Lee and Paul Hongsuck Seo}, year = {2025}, url = {https://arxiv.org/abs/2506.06537}, note = {Source identifier: 2506.06537} }