@misc{indiciaefbd4b28e396f, title = {VLA-JEPA: Enhancing Vision-Language-Action Model with Latent World Model}, author = {Jingwen Sun and Wenyao Zhang and Zekun Qi and Shaojie Ren and Zezhi Liu and Hanxin Zhu and Guangzhong Sun and Xin Jin and Zhibo Chen}, year = {2026}, url = {https://arxiv.org/abs/2602.10098}, note = {Source identifier: 2602.10098} }