@misc{indiciaee49a439e91bb, title = {DetCLIPv2: Scalable Open-Vocabulary Object Detection Pre-training via Word-Region Alignment}, author = {Lewei Yao and Jianhua Han and Xiaodan Liang and Dan Xu and Wei Zhang and Zhenguo Li and Hang Xu}, year = {2023}, url = {https://arxiv.org/abs/2304.04514}, note = {Source identifier: 2304.04514} }