@misc{indiciae6a83c9d03b26, title = {Mantis: A Versatile Vision-Language-Action Model with Disentangled Visual Foresight}, author = {Yi Yang and Xueqi Li and Yiyang Chen and Jin Song and Yihan Wang and Zipeng Xiao and Jiadi Su and You Qiaoben and Pengfei Liu and Zhijie Deng}, year = {2026}, url = {https://arxiv.org/abs/2511.16175}, note = {Source identifier: 2511.16175} }