Views
No views yet

1@misc{podell2023sdxl,
2 title={SDXL: Improving Latent Diffusion Models for High-Resolution Image Synthesis},
3 author={Dustin Podell and Zion English and Kyle Lacey and Andreas Blattmann and Tim Dockhorn and Jonas Müller and Joe Penna and Robin Rombach},
4 year={2023},
5 eprint={2307.01952},
6 archivePrefix={arXiv},
7 primaryClass={cs.CV}
8}pip install invisible_watermark transformers accelerate safetensors diffusers1from diffusers import StableDiffusionXLPipeline, EulerAncestralDiscreteScheduler
2import torch
3
4model_id = "aipicasso/emi-2-5"
5
6scheduler = EulerAncestralDiscreteScheduler.from_pretrained(model_id,subfolder="scheduler")
7pipe = StableDiffusionXLPipeline.from_pretrained(model_id, scheduler=scheduler, torch_dtype=torch.bfloat16)
8pipe = pipe.to("cuda")
9
10prompt = "1girl, upper body, brown bob short hair, brown eyes, looking at viewer, cherry blossom"
11images = pipe(prompt, num_inference_steps=20).images
12images[0].save("girl.png")
131@misc{podell2023sdxl,
2 title={SDXL: Improving Latent Diffusion Models for High-Resolution Image Synthesis},
3 author={Dustin Podell and Zion English and Kyle Lacey and Andreas Blattmann and Tim Dockhorn and Jonas Müller and Joe Penna and Robin Rombach},
4 year={2023},
5 eprint={2307.01952},
6 archivePrefix={arXiv},
7 primaryClass={cs.CV}
8}1@article{li2024cosmicman,
2 title={CosmicMan: A Text-to-Image Foundation Model for Humans},
3 author={Li, Shikai and Fu, Jianglin and Liu, Kaiyuan and Wang, Wentao and Lin, Kwan-Yee and Wu, Wayne},
4 journal={arXiv preprint arXiv:2404.01294},
5 year={2024}
6}
7