hyperparam tweaks

This commit is contained in:
xander
2024-04-26 08:45:29 +02:00
parent af6c4676d5
commit 37bcfe61c9
5 changed files with 48 additions and 36 deletions
+14 -12
View File
@@ -35,11 +35,11 @@ def hamming_distance(dict1, dict2):
#######################################################################################
# Setup the base experiment config:
exp_name = "gridsearch_sdxl_beeple"
exp_name = "gridsearch_sdxl_styles"
caption_prefix = ""
mask_target_prompts = ""
n_exp = 200 # how many random experiment settings to generate
min_hamming_distance = 2 # min_n_params that have to be different from any previous experiment to be scheduled
min_hamming_distance = 3 # min_n_params that have to be different from any previous experiment to be scheduled
output_sh_path = f"gridsearch_configs/{exp_name}.sh"
@@ -50,7 +50,9 @@ hyperparameters = {
"output_dir": [f"lora_models/{exp_name}"],
"sd_model_version": ["sdxl"],
"lora_training_urls": [
"/home/rednax/Documents/datasets/beeple"
"/home/rednax/Documents/datasets/sweep/beeple",
"/home/rednax/Documents/datasets/sweep/clipx_75",
"/home/rednax/Documents/datasets/sweep/speelveld"
],
"concept_mode": ['style'],
"seed": [0],
@@ -58,21 +60,21 @@ hyperparameters = {
"validation_img_size": [[1024, 1024]],
"train_batch_size": [4],
"n_sample_imgs": [6],
"max_train_steps": [600],
"checkpointing_steps": [100],
"max_train_steps": [400],
"checkpointing_steps": [80],
"gradient_accumulation_steps": [1],
"freeze_ti_after_completion_f": [0.4,0.8],
"tok_cov_reg_w": [0.0, 0.001, 0.005, 0.015],
"text_encoder_lora_optimizer": ["adamw", None],
"prodigy_d_coef": [0.5,1.0],
"cond_reg_w": [0.0],
"tok_cond_reg_w": [0.0],
"token_warmup_steps": [100,0],
"freeze_ti_after_completion_f": [0.5,1.0],
"tok_cov_reg_w": [500,1000,2000],
"text_encoder_lora_optimizer": ["adamw"],
"text_encoder_lora_lr": [0.0, 0.5e-4, 2.0e-5],
"unet_prodigy_growth_factor": [1.05, 1.02, 1.01],
"prodigy_d_coef": [1.0],
"n_tokens": [2,3],
"ti_lr": [0.001],
"ti_weight_decay": [0.0005],
"lora_weight_decay": [0.001],
"l1_penalty": [0.1, 0.0],
"off_ratio_power": [0.0, 0.05],
"token_embedding_lr_warmup_steps": [0],
"snr_gamma": [5.0],
"lora_rank": [12,24],