@misc{indiciae9e9aed24cef8, title = {Self-Training Vision Language BERTs with a Unified Conditional Model}, author = {Xiaofeng Yang and Fengmao Lv and Fayao Liu and Guosheng Lin}, year = {2023}, doi = {10.1109/tcsvt.2023.3235704}, url = {https://arxiv.org/abs/2201.02010}, note = {Source identifier: 2201.02010} }