@misc{indiciae560d22373604, title = {Vision-Language Embodiment for Monocular Depth Estimation}, author = {Jinchang Zhang and Guoyu Lu}, year = {2025}, url = {https://arxiv.org/abs/2503.16535}, note = {Source identifier: 2503.16535} }