@misc{indiciae01145ac678a1, title = {OmniJARVIS: Unified Vision-Language-Action Tokenization Enables Open-World Instruction Following Agents}, author = {Zihao Wang and Shaofei Cai and Zhancun Mu and Haowei Lin and Ceyao Zhang and Xuejie Liu and Qing Li and Anji Liu and Xiaojian Ma and Yitao Liang}, year = {2024}, url = {https://arxiv.org/abs/2407.00114}, note = {Source identifier: 2407.00114} }