@misc{indiciae0e4c46b5d3dc, title = {Transform, Contrast and Tell: Coherent Entity-Aware Multi-Image Captioning}, author = {Jingqiang Chen}, year = {2023}, doi = {10.1016/j.cviu.2023.103878}, url = {https://arxiv.org/abs/2302.02124}, note = {Source identifier: 2302.02124} }