stable-diffusion-webui/modules/textual_inversion/preprocess.py

import os
from PIL import Image, ImageOps
import math
import tqdm

from modules import paths, shared, images, deepbooru
from modules.textual_inversion import autocrop


def preprocess(id_task, process_src, process_dst, process_width, process_height, preprocess_txt_action, process_keep_original_size, process_flip, process_split, process_caption, process_caption_deepbooru=False, split_threshold=0.5, overlap_ratio=0.2, process_focal_crop=False, process_focal_crop_face_weight=0.9, process_focal_crop_entropy_weight=0.3, process_focal_crop_edges_weight=0.5, process_focal_crop_debug=False, process_multicrop=None, process_multicrop_mindim=None, process_multicrop_maxdim=None, process_multicrop_minarea=None, process_multicrop_maxarea=None, process_multicrop_objective=None, process_multicrop_threshold=None):
    try:
        if process_caption:
            shared.interrogator.load()

        if process_caption_deepbooru:
            deepbooru.model.start()

        preprocess_work(process_src, process_dst, process_width, process_height, preprocess_txt_action, process_keep_original_size, process_flip, process_split, process_caption, process_caption_deepbooru, split_threshold, overlap_ratio, process_focal_crop, process_focal_crop_face_weight, process_focal_crop_entropy_weight, process_focal_crop_edges_weight, process_focal_crop_debug, process_multicrop, process_multicrop_mindim, process_multicrop_maxdim, process_multicrop_minarea, process_multicrop_maxarea, process_multicrop_objective, process_multicrop_threshold)

    finally:

        if process_caption:
            shared.interrogator.send_blip_to_ram()

        if process_caption_deepbooru:
            deepbooru.model.stop()


def listfiles(dirname):
    return os.listdir(dirname)


class PreprocessParams:
    src = None
    dstdir = None
    subindex = 0
    flip = False
    process_caption = False
    process_caption_deepbooru = False
    preprocess_txt_action = None


def save_pic_with_caption(image, index, params: PreprocessParams, existing_caption=None):
    caption = ""

    if params.process_caption:
        caption += shared.interrogator.generate_caption(image)

    if params.process_caption_deepbooru:
        if len(caption) > 0:
            caption += ", "
        caption += deepbooru.model.tag_multi(image)

    filename_part = params.src
    filename_part = os.path.splitext(filename_part)[0]
    filename_part = os.path.basename(filename_part)

    basename = f"{index:05}-{params.subindex}-{filename_part}"
    image.save(os.path.join(params.dstdir, f"{basename}.png"))

    if params.preprocess_txt_action == 'prepend' and existing_caption:
        caption = f"{existing_caption} {caption}"
    elif params.preprocess_txt_action == 'append' and existing_caption:
        caption = f"{caption} {existing_caption}"
    elif params.preprocess_txt_action == 'copy' and existing_caption:
        caption = existing_caption

    caption = caption.strip()

    if len(caption) > 0:
        with open(os.path.join(params.dstdir, f"{basename}.txt"), "w", encoding="utf8") as file:
            file.write(caption)

    params.subindex += 1


def save_pic(image, index, params, existing_caption=None):
    save_pic_with_caption(image, index, params, existing_caption=existing_caption)

    if params.flip:
        save_pic_with_caption(ImageOps.mirror(image), index, params, existing_caption=existing_caption)


def split_pic(image, inverse_xy, width, height, overlap_ratio):
    if inverse_xy:
        from_w, from_h = image.height, image.width
        to_w, to_h = height, width
    else:
        from_w, from_h = image.width, image.height
        to_w, to_h = width, height
    h = from_h * to_w // from_w
    if inverse_xy:
        image = image.resize((h, to_w))
    else:
        image = image.resize((to_w, h))

    split_count = math.ceil((h - to_h * overlap_ratio) / (to_h * (1.0 - overlap_ratio)))
    y_step = (h - to_h) / (split_count - 1)
    for i in range(split_count):
        y = int(y_step * i)
        if inverse_xy:
            splitted = image.crop((y, 0, y + to_h, to_w))
        else:
            splitted = image.crop((0, y, to_w, y + to_h))
        yield splitted

# not using torchvision.transforms.CenterCrop because it doesn't allow float regions
def center_crop(image: Image, w: int, h: int):
    iw, ih = image.size
    if ih / h < iw / w:
        sw = w * ih / h
        box = (iw - sw) / 2, 0, iw - (iw - sw) / 2, ih
    else:
        sh = h * iw / w
        box = 0, (ih - sh) / 2, iw, ih - (ih - sh) / 2
    return image.resize((w, h), Image.Resampling.LANCZOS, box)


def multicrop_pic(image: Image, mindim, maxdim, minarea, maxarea, objective, threshold):
    iw, ih = image.size
    err = lambda w, h: 1-(lambda x: x if x < 1 else 1/x)(iw/ih/(w/h))
    wh = max(((w, h) for w in range(mindim, maxdim+1, 64) for h in range(mindim, maxdim+1, 64)
        if minarea <= w * h <= maxarea and err(w, h) <= threshold),
        key= lambda wh: (wh[0]*wh[1], -err(*wh))[::1 if objective=='Maximize area' else -1],
        default=None
    )
    return wh and center_crop(image, *wh)
    

def preprocess_work(process_src, process_dst, process_width, process_height, preprocess_txt_action, process_keep_original_size, process_flip, process_split, process_caption, process_caption_deepbooru=False, split_threshold=0.5, overlap_ratio=0.2, process_focal_crop=False, process_focal_crop_face_weight=0.9, process_focal_crop_entropy_weight=0.3, process_focal_crop_edges_weight=0.5, process_focal_crop_debug=False, process_multicrop=None, process_multicrop_mindim=None, process_multicrop_maxdim=None, process_multicrop_minarea=None, process_multicrop_maxarea=None, process_multicrop_objective=None, process_multicrop_threshold=None):
    width = process_width
    height = process_height
    src = os.path.abspath(process_src)
    dst = os.path.abspath(process_dst)
    split_threshold = max(0.0, min(1.0, split_threshold))
    overlap_ratio = max(0.0, min(0.9, overlap_ratio))

    assert src != dst, 'same directory specified as source and destination'

    os.makedirs(dst, exist_ok=True)

    files = listfiles(src)

    shared.state.job = "preprocess"
    shared.state.textinfo = "Preprocessing..."
    shared.state.job_count = len(files)

    params = PreprocessParams()
    params.dstdir = dst
    params.flip = process_flip
    params.process_caption = process_caption
    params.process_caption_deepbooru = process_caption_deepbooru
    params.preprocess_txt_action = preprocess_txt_action

    pbar = tqdm.tqdm(files)
    for index, imagefile in enumerate(pbar):
        params.subindex = 0
        filename = os.path.join(src, imagefile)
        try:
            img = Image.open(filename)
            img = ImageOps.exif_transpose(img)
            img = img.convert("RGB")
        except Exception:
            continue

        description = f"Preprocessing [Image {index}/{len(files)}]"
        pbar.set_description(description)
        shared.state.textinfo = description

        params.src = filename

        existing_caption = None
        existing_caption_filename = f"{os.path.splitext(filename)[0]}.txt"
        if os.path.exists(existing_caption_filename):
            with open(existing_caption_filename, 'r', encoding="utf8") as file:
                existing_caption = file.read()

        if shared.state.interrupted:
            break

        if img.height > img.width:
            ratio = (img.width * height) / (img.height * width)
            inverse_xy = False
        else:
            ratio = (img.height * width) / (img.width * height)
            inverse_xy = True

        process_default_resize = True

        if process_split and ratio < 1.0 and ratio <= split_threshold:
            for splitted in split_pic(img, inverse_xy, width, height, overlap_ratio):
                save_pic(splitted, index, params, existing_caption=existing_caption)
            process_default_resize = False

        if process_focal_crop and img.height != img.width:

            dnn_model_path = None
            try:
                dnn_model_path = autocrop.download_and_cache_models(os.path.join(paths.models_path, "opencv"))
            except Exception as e:
                print("Unable to load face detection model for auto crop selection. Falling back to lower quality haar method.", e)

            autocrop_settings = autocrop.Settings(
                crop_width = width,
                crop_height = height,
                face_points_weight = process_focal_crop_face_weight,
                entropy_points_weight = process_focal_crop_entropy_weight,
                corner_points_weight = process_focal_crop_edges_weight,
                annotate_image = process_focal_crop_debug,
                dnn_model_path = dnn_model_path,
            )
            for focal in autocrop.crop_image(img, autocrop_settings):
                save_pic(focal, index, params, existing_caption=existing_caption)
            process_default_resize = False

        if process_multicrop:
            cropped = multicrop_pic(img, process_multicrop_mindim, process_multicrop_maxdim, process_multicrop_minarea, process_multicrop_maxarea, process_multicrop_objective, process_multicrop_threshold)
            if cropped is not None:
                save_pic(cropped, index, params, existing_caption=existing_caption)
            else:
                print(f"skipped {img.width}x{img.height} image {filename} (can't find suitable size within error threshold)")
            process_default_resize = False

        if process_keep_original_size:
            save_pic(img, index, params, existing_caption=existing_caption)
            process_default_resize = False

        if process_default_resize:
            img = images.resize_image(1, img, width, height)
            save_pic(img, index, params, existing_caption=existing_caption)

        shared.state.nextjob()
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00			`import os`
face detection algo, configurability, reusability Try to move the crop in the direction of a face if it is present More internal configuration options for choosing weights of each of the algorithm's findings Move logic into its module 2022-10-20 03:19:02 +03:00			`from PIL import Image, ImageOps`
train: fixed preprocess image ratio 2022-10-20 10:53:46 +03:00			`import math`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00			`import tqdm`

add data-dir flag and set all user data directories based on it 2023-01-25 19:15:42 +03:00			`from modules import paths, shared, images, deepbooru`
face detection algo, configurability, reusability Try to move the crop in the direction of a face if it is present More internal configuration options for choosing weights of each of the algorithm's findings Move logic into its module 2022-10-20 03:19:02 +03:00			`from modules.textual_inversion import autocrop`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
deepbooru: added option to use spaces or underscores deepbooru: added option to quote (\) in tags deepbooru/BLIP: write caption to file instead of image filename deepbooru/BLIP: now possible to use both for captions deepbooru: process is stopped even if an exception occurs 2022-10-12 21:55:43 +03:00
Add option "keep original size" to textual inversion images preprocess 2023-03-25 17:45:41 +03:00			def preprocess(id_task, process_src, process_dst, process_width, process_height, preprocess_txt_action, process_keep_original_size, process_flip, process_split, process_caption, process_caption_deepbooru=False, split_threshold=0.5, overlap_ratio=0.2, process_focal_crop=False, process_focal_crop_face_weight=0.9, process_focal_crop_entropy_weight=0.3, process_focal_crop_edges_weight=0.5, process_focal_crop_debug=False, process_multicrop=None, process_multicrop_mindim=None, process_multicrop_maxdim=None, process_multicrop_minarea=None, process_multicrop_maxarea=None, process_multicrop_objective=None, process_multicrop_threshold=None):
deepbooru: added option to use spaces or underscores deepbooru: added option to quote (\) in tags deepbooru/BLIP: write caption to file instead of image filename deepbooru/BLIP: now possible to use both for captions deepbooru: process is stopped even if an exception occurs 2022-10-12 21:55:43 +03:00			`try:`
			`if process_caption:`
			`shared.interrogator.load()`

			`if process_caption_deepbooru:`
moved deepdanbooru to pure pytorch implementation 2022-11-20 16:39:20 +03:00			`deepbooru.model.start()`
deepbooru: added option to use spaces or underscores deepbooru: added option to quote (\) in tags deepbooru/BLIP: write caption to file instead of image filename deepbooru/BLIP: now possible to use both for captions deepbooru: process is stopped even if an exception occurs 2022-10-12 21:55:43 +03:00
Add option "keep original size" to textual inversion images preprocess 2023-03-25 17:45:41 +03:00			preprocess_work(process_src, process_dst, process_width, process_height, preprocess_txt_action, process_keep_original_size, process_flip, process_split, process_caption, process_caption_deepbooru, split_threshold, overlap_ratio, process_focal_crop, process_focal_crop_face_weight, process_focal_crop_entropy_weight, process_focal_crop_edges_weight, process_focal_crop_debug, process_multicrop, process_multicrop_mindim, process_multicrop_maxdim, process_multicrop_minarea, process_multicrop_maxarea, process_multicrop_objective, process_multicrop_threshold)
deepbooru: added option to use spaces or underscores deepbooru: added option to quote (\) in tags deepbooru/BLIP: write caption to file instead of image filename deepbooru/BLIP: now possible to use both for captions deepbooru: process is stopped even if an exception occurs 2022-10-12 21:55:43 +03:00
			`finally:`

			`if process_caption:`
			`shared.interrogator.send_blip_to_ram()`

			`if process_caption_deepbooru:`
moved deepdanbooru to pure pytorch implementation 2022-11-20 16:39:20 +03:00			`deepbooru.model.stop()`
deepbooru: added option to use spaces or underscores deepbooru: added option to quote (\) in tags deepbooru/BLIP: write caption to file instead of image filename deepbooru/BLIP: now possible to use both for captions deepbooru: process is stopped even if an exception occurs 2022-10-12 21:55:43 +03:00

move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`def listfiles(dirname):`
			`return os.listdir(dirname)`


			`class PreprocessParams:`
			`src = None`
			`dstdir = None`
			`subindex = 0`
			`flip = False`
			`process_caption = False`
			`process_caption_deepbooru = False`
			`preprocess_txt_action = None`


			`def save_pic_with_caption(image, index, params: PreprocessParams, existing_caption=None):`
			`caption = ""`

			`if params.process_caption:`
			`caption += shared.interrogator.generate_caption(image)`

			`if params.process_caption_deepbooru:`
			`if len(caption) > 0:`
			`caption += ", "`
moved deepdanbooru to pure pytorch implementation 2022-11-20 16:39:20 +03:00			`caption += deepbooru.model.tag_multi(image)`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00
			`filename_part = params.src`
			`filename_part = os.path.splitext(filename_part)[0]`
			`filename_part = os.path.basename(filename_part)`

			`basename = f"{index:05}-{params.subindex}-{filename_part}"`
			`image.save(os.path.join(params.dstdir, f"{basename}.png"))`

			`if params.preprocess_txt_action == 'prepend' and existing_caption:`
Fix up string formatting/concatenation to f-strings where feasible 2023-05-09 22:17:58 +03:00			`caption = f"{existing_caption} {caption}"`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`elif params.preprocess_txt_action == 'append' and existing_caption:`
Fix up string formatting/concatenation to f-strings where feasible 2023-05-09 22:17:58 +03:00			`caption = f"{caption} {existing_caption}"`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`elif params.preprocess_txt_action == 'copy' and existing_caption:`
			`caption = existing_caption`

			`caption = caption.strip()`

			`if len(caption) > 0:`
			`with open(os.path.join(params.dstdir, f"{basename}.txt"), "w", encoding="utf8") as file:`
			`file.write(caption)`

			`params.subindex += 1`


			`def save_pic(image, index, params, existing_caption=None):`
			`save_pic_with_caption(image, index, params, existing_caption=existing_caption)`

			`if params.flip:`
			`save_pic_with_caption(ImageOps.mirror(image), index, params, existing_caption=existing_caption)`


			`def split_pic(image, inverse_xy, width, height, overlap_ratio):`
			`if inverse_xy:`
			`from_w, from_h = image.height, image.width`
			`to_w, to_h = height, width`
			`else:`
			`from_w, from_h = image.width, image.height`
			`to_w, to_h = width, height`
			`h = from_h * to_w // from_w`
			`if inverse_xy:`
			`image = image.resize((h, to_w))`
			`else:`
			`image = image.resize((to_w, h))`

			`split_count = math.ceil((h - to_h * overlap_ratio) / (to_h * (1.0 - overlap_ratio)))`
			`y_step = (h - to_h) / (split_count - 1)`
			`for i in range(split_count):`
			`y = int(y_step * i)`
			`if inverse_xy:`
			`splitted = image.crop((y, 0, y + to_h, to_w))`
			`else:`
			`splitted = image.crop((0, y, to_w, y + to_h))`
			`yield splitted`

Add auto-sized cropping UI 2023-01-17 12:16:43 +03:00			`# not using torchvision.transforms.CenterCrop because it doesn't allow float regions`
			`def center_crop(image: Image, w: int, h: int):`
			`iw, ih = image.size`
			`if ih / h < iw / w:`
			`sw = w * ih / h`
			`box = (iw - sw) / 2, 0, iw - (iw - sw) / 2, ih`
			`else:`
			`sh = h * iw / w`
			`box = 0, (ih - sh) / 2, iw, ih - (ih - sh) / 2`
			`return image.resize((w, h), Image.Resampling.LANCZOS, box)`

deepbooru: added option to use spaces or underscores deepbooru: added option to quote (\) in tags deepbooru/BLIP: write caption to file instead of image filename deepbooru/BLIP: now possible to use both for captions deepbooru: process is stopped even if an exception occurs 2022-10-12 21:55:43 +03:00
Add auto-sized cropping UI 2023-01-17 12:16:43 +03:00			`def multicrop_pic(image: Image, mindim, maxdim, minarea, maxarea, objective, threshold):`
			`iw, ih = image.size`
			`err = lambda w, h: 1-(lambda x: x if x < 1 else 1/x)(iw/ih/(w/h))`
Fix of fix 2023-01-19 12:39:30 +03:00			`wh = max(((w, h) for w in range(mindim, maxdim+1, 64) for h in range(mindim, maxdim+1, 64)`
Simplification and bugfix 2023-01-19 12:36:23 +03:00			`if minarea <= w * h <= maxarea and err(w, h) <= threshold),`
			`key= lambda wh: (wh[0]wh[1], -err(wh))[::1 if objective=='Maximize area' else -1],`
			`default=None`
			`)`
Fix of fix 2023-01-19 12:39:30 +03:00			`return wh and center_crop(image, *wh)`
Add auto-sized cropping UI 2023-01-17 12:16:43 +03:00

Add option "keep original size" to textual inversion images preprocess 2023-03-25 17:45:41 +03:00			def preprocess_work(process_src, process_dst, process_width, process_height, preprocess_txt_action, process_keep_original_size, process_flip, process_split, process_caption, process_caption_deepbooru=False, split_threshold=0.5, overlap_ratio=0.2, process_focal_crop=False, process_focal_crop_face_weight=0.9, process_focal_crop_entropy_weight=0.3, process_focal_crop_edges_weight=0.5, process_focal_crop_debug=False, process_multicrop=None, process_multicrop_mindim=None, process_multicrop_maxdim=None, process_multicrop_minarea=None, process_multicrop_maxarea=None, process_multicrop_objective=None, process_multicrop_threshold=None):
Custom Width and Height 2022-10-10 16:35:35 +03:00			`width = process_width`
			`height = process_height`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00			`src = os.path.abspath(process_src)`
			`dst = os.path.abspath(process_dst)`
train: ui: added `Split image threshold` and `Split image overlap ratio` to preprocess 2022-10-20 16:56:45 +03:00			`split_threshold = max(0.0, min(1.0, split_threshold))`
			`overlap_ratio = max(0.0, min(0.9, overlap_ratio))`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
removed unused import, fixed typo 2022-10-06 00:11:32 +03:00			`assert src != dst, 'same directory specified as source and destination'`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
			`os.makedirs(dst, exist_ok=True)`

move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`files = listfiles(src)`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
add job info to modules 2023-01-03 18:34:51 +03:00			`shared.state.job = "preprocess"`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00			`shared.state.textinfo = "Preprocessing..."`
			`shared.state.job_count = len(files)`

move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`params = PreprocessParams()`
			`params.dstdir = dst`
			`params.flip = process_flip`
			`params.process_caption = process_caption`
			`params.process_caption_deepbooru = process_caption_deepbooru`
			`params.preprocess_txt_action = preprocess_txt_action`
face detection algo, configurability, reusability Try to move the crop in the direction of a face if it is present More internal configuration options for choosing weights of each of the algorithm's findings Move logic into its module 2022-10-20 03:19:02 +03:00
set descriptions 2023-01-11 18:28:55 +03:00			`pbar = tqdm.tqdm(files)`
			`for index, imagefile in enumerate(pbar):`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`params.subindex = 0`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00			`filename = os.path.join(src, imagefile)`
Switched to exception handling 2022-10-11 11:32:46 +03:00			`try:`
fix preprocess orientation 2023-04-06 02:28:00 +03:00			`img = Image.open(filename)`
Pythonic way to achieve it 2023-04-06 02:51:29 +03:00			`img = ImageOps.exif_transpose(img)`
fix preprocess orientation 2023-04-06 02:28:00 +03:00			`img = img.convert("RGB")`
Switched to exception handling 2022-10-11 11:32:46 +03:00			`except Exception:`
			`continue`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
set descriptions 2023-01-11 18:28:55 +03:00			`description = f"Preprocessing [Image {index}/{len(files)}]"`
			`pbar.set_description(description)`
			`shared.state.textinfo = description`

move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`params.src = filename`

add existing caption file handling 2022-10-20 02:46:54 +03:00			`existing_caption = None`
Fix up string formatting/concatenation to f-strings where feasible 2023-05-09 22:17:58 +03:00			`existing_caption_filename = f"{os.path.splitext(filename)[0]}.txt"`
prevent error spam when processing images without txt files for captions 2022-10-21 18:46:02 +03:00			`if os.path.exists(existing_caption_filename):`
			`with open(existing_caption_filename, 'r', encoding="utf8") as file:`
			`existing_caption = file.read()`
add existing caption file handling 2022-10-20 02:46:54 +03:00
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00			`if shared.state.interrupted:`
			`break`

train: fixed preprocess image ratio 2022-10-20 10:53:46 +03:00			`if img.height > img.width:`
			`ratio = (img.width * height) / (img.height * width)`
			`inverse_xy = False`
			`else:`
			`ratio = (img.height * width) / (img.width * height)`
			`inverse_xy = True`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
Focal crop UI elements 2022-10-26 01:22:29 +03:00			`process_default_resize = True`
Add auto focal point cropping to Preprocess images This algorithm plots a bunch of points of interest on the source image and averages their locations to find a center. Most points come from OpenCV. One point comes from an entropy model. OpenCV points account for 50% of the weight and the entropy based point is the other 50%. The center of all weighted points is calculated and a bounding box is drawn as close to centered over that point as possible. 2022-10-19 13:18:26 +03:00
train: fixed preprocess image ratio 2022-10-20 10:53:46 +03:00			`if process_split and ratio < 1.0 and ratio <= split_threshold:`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`for splitted in split_pic(img, inverse_xy, width, height, overlap_ratio):`
			`save_pic(splitted, index, params, existing_caption=existing_caption)`
Focal crop UI elements 2022-10-26 01:22:29 +03:00			`process_default_resize = False`
Add auto focal point cropping to Preprocess images This algorithm plots a bunch of points of interest on the source image and averages their locations to find a center. Most points come from OpenCV. One point comes from an entropy model. OpenCV points account for 50% of the weight and the entropy based point is the other 50%. The center of all weighted points is calculated and a bounding box is drawn as close to centered over that point as possible. 2022-10-19 13:18:26 +03:00
download better face detection module dynamically 2022-10-26 02:14:13 +03:00			`if process_focal_crop and img.height != img.width:`

			`dnn_model_path = None`
			`try:`
add data-dir flag and set all user data directories based on it 2023-01-25 19:15:42 +03:00			`dnn_model_path = autocrop.download_and_cache_models(os.path.join(paths.models_path, "opencv"))`
download better face detection module dynamically 2022-10-26 02:14:13 +03:00			`except Exception as e:`
			`print("Unable to load face detection model for auto crop selection. Falling back to lower quality haar method.", e)`

face detection algo, configurability, reusability Try to move the crop in the direction of a face if it is present More internal configuration options for choosing weights of each of the algorithm's findings Move logic into its module 2022-10-20 03:19:02 +03:00			`autocrop_settings = autocrop.Settings(`
			`crop_width = width,`
			`crop_height = height,`
Focal crop UI elements 2022-10-26 01:22:29 +03:00			`face_points_weight = process_focal_crop_face_weight,`
			`entropy_points_weight = process_focal_crop_entropy_weight,`
			`corner_points_weight = process_focal_crop_edges_weight,`
download better face detection module dynamically 2022-10-26 02:14:13 +03:00			`annotate_image = process_focal_crop_debug,`
			`dnn_model_path = dnn_model_path,`
face detection algo, configurability, reusability Try to move the crop in the direction of a face if it is present More internal configuration options for choosing weights of each of the algorithm's findings Move logic into its module 2022-10-20 03:19:02 +03:00			`)`
Focal crop UI elements 2022-10-26 01:22:29 +03:00			`for focal in autocrop.crop_image(img, autocrop_settings):`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`save_pic(focal, index, params, existing_caption=existing_caption)`
Focal crop UI elements 2022-10-26 01:22:29 +03:00			`process_default_resize = False`
Add auto focal point cropping to Preprocess images This algorithm plots a bunch of points of interest on the source image and averages their locations to find a center. Most points come from OpenCV. One point comes from an entropy model. OpenCV points account for 50% of the weight and the entropy based point is the other 50%. The center of all weighted points is calculated and a bounding box is drawn as close to centered over that point as possible. 2022-10-19 13:18:26 +03:00
Add auto-sized cropping UI 2023-01-17 12:16:43 +03:00			`if process_multicrop:`
			`cropped = multicrop_pic(img, process_multicrop_mindim, process_multicrop_maxdim, process_multicrop_minarea, process_multicrop_maxarea, process_multicrop_objective, process_multicrop_threshold)`
			`if cropped is not None:`
			`save_pic(cropped, index, params, existing_caption=existing_caption)`
			`else:`
			`print(f"skipped {img.width}x{img.height} image {filename} (can't find suitable size within error threshold)")`
			`process_default_resize = False`

Add option "keep original size" to textual inversion images preprocess 2023-03-25 17:45:41 +03:00			`if process_keep_original_size:`
			`save_pic(img, index, params, existing_caption=existing_caption)`
			`process_default_resize = False`

Focal crop UI elements 2022-10-26 01:22:29 +03:00			`if process_default_resize:`
Custom Width and Height 2022-10-10 16:35:35 +03:00			`img = images.resize_image(1, img, width, height)`
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`save_pic(img, index, params, existing_caption=existing_caption)`
preprocessing for textual inversion added 2022-10-02 22:41:21 +03:00
move functions out of main body for image preprocessing for easier hijacking 2022-11-08 08:37:05 +03:00			`shared.state.nextjob()`