@misc{indiciaeb16c387d1933, title = {HiFi-CS: Towards Open Vocabulary Visual Grounding For Robotic Grasping Using Vision-Language Models}, author = {Vineet Bhat and Prashanth Krishnamurthy and Ramesh Karri and Farshad Khorrami}, year = {2025}, url = {https://arxiv.org/abs/2409.10419}, note = {Source identifier: 2409.10419} }