@misc{indiciae959236310e27, title = {OmDet: Large-scale vision-language multi-dataset pre-training with multimodal detection network}, author = {Tiancheng Zhao and Peng Liu and Kyusong Lee}, year = {2024}, doi = {10.1049/cvi2.12268}, url = {https://arxiv.org/abs/2209.05946}, note = {Source identifier: 2209.05946} }