@misc{indiciae0bbcc4c93969, title = {Multi-modal Understanding and Generation for Medical Images and Text via Vision-Language Pre-Training}, author = {Jong Hak Moon and Hyungyung Lee and Woncheol Shin and Young-Hak Kim and Edward Choi}, year = {2022}, doi = {10.1109/jbhi.2022.3207502}, url = {https://arxiv.org/abs/2105.11333}, note = {Source identifier: 2105.11333} }