@misc{indiciaeeef353306e58, title = {Constructing Multi-Modal Dialogue Dataset by Replacing Text with Semantically Relevant Images}, author = {Nyoungwoo Lee and Suwon Shin and Jaegul Choo and Ho-Jin Choi and Sung-Hyun Myaeng}, year = {2021}, url = {https://arxiv.org/abs/2107.08685}, note = {Source identifier: 2107.08685} }