@misc{indiciae9f7248e186ae, title = {SAMOT: State-Aware Step Modulation and Optimal Transport Matching for Audio-Visual Instance Segmentation}, author = {Kai Peng and Yunzhe Shen and Miao Zhang and Leiye Liu and Wei Ji and Jingjing Li and Yongri Piao and Huchuan Lu}, year = {2026}, url = {https://arxiv.org/abs/2608.08147}, note = {Source identifier: 2608.08147} }