@misc{indiciae024cb66b3490, title = {OPERA: A Unified Omnimodal Progressive Spatio-Temporal Reasoning Agent for Referring Video Segmentation}, author = {Jingchen Ni and Yuji Wang and Shannan Yan and Haoru Li and Sitong Chen and Chun Yuan}, year = {2026}, url = {https://arxiv.org/abs/2609.33338}, note = {Source identifier: 2609.33338} }