@misc{indiciae564133278fd1, title = {ChatVLA-2: Vision-Language-Action Model with Open-World Embodied Reasoning from Pretrained Knowledge}, author = {Zhongyi Zhou and Yichen Zhu and Junjie Wen and Chaomin Shen and Yi Xu}, year = {2025}, url = {https://arxiv.org/abs/2505.21906}, note = {Source identifier: 2505.21906} }