@misc{indiciae6aad3ed5198f, title = {Enhancing Open-Vocabulary Object Detection through Multi-Level Fine-Grained Visual-Language Alignment}, author = {Tianyi Zhang and Antoine Simoulin and Kai Li and Sana Lakdawala and Shiqing Yu and Arpit Mittal and Hongyu Fu and Yu Lin}, year = {2026}, url = {https://arxiv.org/abs/2602.00531}, note = {Source identifier: 2602.00531} }