@misc{indiciaea5126470458c, title = {VIRST: Video-Instructed Reasoning Assistant for SpatioTemporal Segmentation}, author = {Jihwan Hong and Jaeyoung Do}, year = {2026}, url = {https://arxiv.org/abs/2603.27060}, note = {Source identifier: 2603.27060} }