@misc{indiciaeb306352c43ce, title = {e-ViL: A Dataset and Benchmark for Natural Language Explanations in Vision-Language Tasks}, author = {Maxime Kayser and Oana-Maria Camburu and Leonard Salewski and Cornelius Emde and Virginie Do and Zeynep Akata and Thomas Lukasiewicz}, year = {2021}, url = {https://arxiv.org/abs/2105.03761}, note = {Source identifier: 2105.03761} }