@inproceedings{7d96fd583d944ba4b62e098e40d2d1f0,
title = "LOOSECONTROL: Lifting ControlNet for Generalized Depth Conditioning",
abstract = "We present LooseControl to allow generalized depth conditioning for diffusion-based image generation. ControlNet, the SOTA for depth conditioned image generation, produces remarkable results but relies on having access to detailed depth maps for guidance. Creating such exact depth maps, in many scenarios, is challenging. This paper introduces a generalized version of depth conditioning that enables new content creation workflows. Specifically, we allow (C1) scene boundary control for loosely specifying scenes with only boundary conditions, and (C2) 3D box control for specifying the target objects' layout locations rather than the objects' exact shape and appearance. Using LooseControl, along with text guidance, users can create complex environments (e.g., rooms, street views, etc.) by specifying only scene boundaries and locations of primary objects. Further, we provide two editing mechanisms to refine the results: (E1) 3D box editing enables the user to refine images by changing, adding, or removing boxes while freezing the image style. This yields minimal changes apart from changes induced by the edited boxes. (E2) Attribute editing proposes possible editing directions to change one particular aspect of the scene, such as the overall object density or a particular object. Tests and comparisons with baselines demonstrate the generality of our method. We believe that LooseControl can become an important design tool for easily creating complex environments and be extended to other forms of guidance channels.",
keywords = "control, depth condition, diffusion models, generative models, guided editing, layout control, partial specification",
author = "Bhat, \{Shariq Farooq\} and Niloy Mitra and Peter Wonka",
note = "Publisher Copyright: {\textcopyright} 2024 Owner/Author.; 2024 Special Interest Group on Computer Graphics and Interactive Techniques Conference - Conference Papers, SIGGRAPH 2024 ; Conference date: 28-07-2024 Through 01-08-2024",
year = "2024",
month = jul,
day = "13",
doi = "10.1145/3641519.3657525",
language = "English (US)",
series = "Proceedings - SIGGRAPH 2024 Conference Papers",
publisher = "Association for Computing Machinery, Inc",
editor = "Spencer, \{Stephen N.\}",
booktitle = "Proceedings - SIGGRAPH 2024 Conference Papers",
}