@misc{indiciae031ed2703f88, title = {A Training-Free Guess What Vision Language Model from Snippets to Open-Vocabulary Object Detection}, author = {Guiying Zhu and Bowen Yang and Yin Zhuang and Tong Zhang and Guanqun Wang and Zhihao Che and He Chen and Lianlin Li}, year = {2026}, url = {https://arxiv.org/abs/2601.11910}, note = {Source identifier: 2601.11910} }