@misc{indiciae492c82755d91, title = {VELMA: Verbalization Embodiment of LLM Agents for Vision and Language Navigation in Street View}, author = {Raphael Schumann and Wanrong Zhu and Weixi Feng and Tsu-Jui Fu and Stefan Riezler and William Yang Wang}, year = {2024}, url = {https://arxiv.org/abs/2307.06082}, note = {Source identifier: 2307.06082} }