@misc{indiciaed13c90ecaab3, title = {Do Vision-Language-Action Models Understand Instructions? A Mechanistic Interpretability Study on Language Grounding}, author = {Theodor Wulff and Angelo Cangelosi}, year = {2026}, url = {https://arxiv.org/abs/2610.10178}, note = {Source identifier: 2610.10178} }