@misc{indiciaef212f7b17622, title = {EmbodimentSemantic: A Spatial Scene-Graph Dataset and Benchmark for Vision-Language Models on Embodied Manipulation Trajectories}, author = {Hassan Jaber and Refinath S N and Luca Cagliero and Christopher E. Mower and Haitham Bou-Ammar}, year = {2026}, url = {https://arxiv.org/abs/2607.00020}, note = {Source identifier: 2607.00020} }