@misc{indiciae85b1a23aca3e, title = {EmbSpatial-Bench: Benchmarking Spatial Understanding for Embodied Tasks with Large Vision-Language Models}, author = {Mengfei Du and Binhao Wu and Zejun Li and Xuanjing Huang and Zhongyu Wei}, year = {2024}, url = {https://arxiv.org/abs/2406.05756}, note = {Source identifier: 2406.05756} }