@misc{indiciae8af40762efcd, title = {Physically Grounded Vision-Language Models for Robotic Manipulation}, author = {Jensen Gao and Bidipta Sarkar and Fei Xia and Ted Xiao and Jiajun Wu and Brian Ichter and Anirudha Majumdar and Dorsa Sadigh}, year = {2024}, url = {https://arxiv.org/abs/2309.02561}, note = {Source identifier: 2309.02561} }