@misc{indiciae23d4813c9c2f, title = {More Grounded Image Captioning by Distilling Image-Text Matching Model}, author = {Yuanen Zhou and Meng Wang and Daqing Liu and Zhenzhen Hu and Hanwang Zhang}, year = {2020}, url = {https://arxiv.org/abs/2004.00390}, note = {Source identifier: 2004.00390} }