@misc{indiciae24673f87cd9c, title = {ManipBench: Benchmarking Vision-Language Models for Low-Level Robot Manipulation}, author = {Enyu Zhao and Vedant Raval and Hejia Zhang and Jiageng Mao and Zeyu Shangguan and Stefanos Nikolaidis and Yue Wang and Daniel Seita}, year = {2025}, url = {https://arxiv.org/abs/2505.09698}, note = {Source identifier: 2505.09698} }