@misc{indiciae2b072a762736, title = {Conceptual 12M: Pushing Web-Scale Image-Text Pre-Training To Recognize Long-Tail Visual Concepts}, author = {Soravit Changpinyo and Piyush Sharma and Nan Ding and Radu Soricut}, year = {2021}, url = {https://arxiv.org/abs/2102.08981}, note = {Source identifier: 2102.08981} }