@misc{indiciae7a7aa9212e35, title = {Cocktail: Mixing Multi-Modality Controls for Text-Conditional Image Generation}, author = {Minghui Hu and Jianbin Zheng and Daqing Liu and Chuanxia Zheng and Chaoyue Wang and Dacheng Tao and Tat-Jen Cham}, year = {2023}, url = {https://arxiv.org/abs/2306.00964}, note = {Source identifier: 2306.00964} }