@misc{indiciaeefcd81a0cf59, title = {Evo-0: Vision-Language-Action Model with Implicit Spatial Understanding}, author = {Tao Lin and Gen Li and Yilei Zhong and Yanwen Zou and Yuxin Du and Jiting Liu and Encheng Gu and Bo Zhao}, year = {2025}, url = {https://arxiv.org/abs/2507.00416}, note = {Source identifier: 2507.00416} }