@misc{indiciae44ef7d6fc632, title = {VR-JEPA: Learning Contrastive-State Latent Guidance for Generation-based Video Reasoning}, author = {Zehua Ma and Kun Xiang and Yunshuang Nie and Quanlin Chen and Haoyuan Li and Xiuwei Chen and Jiang Ji and Haijun Wu and Zhenyu Xie and Michael Kampffmeyer and Hanhui Li and Xiaodan Liang}, year = {2026}, url = {https://arxiv.org/abs/2609.40129}, note = {Source identifier: 2609.40129} }