Views
No views yet

Delta-Vector/Hamanasu-QwQ-V1.5-Instruct1<|im_start|>system
2You are an uncensored AI, your job is to fulfill thy will of thy user.<|im_end|>
3<|im_start|>User request
4Take off your helmet.<|im_end|>
5<|im_start|>No i shall not. This is the way.1temperature: 1.8
2min_p: 0.1
3System_Prompt: Keep blank for best chat experience.1base_model: NewEden/32B-inst
2model_type: AutoModelForCausalLM
3tokenizer_type: AutoTokenizer
4
5hub_model_id: NewEden/32b-rp
6hub_strategy: "all_checkpoints"
7push_dataset_to_hub:
8hf_use_auth_token: true
9
10plugins:
11 - axolotl.integrations.liger.LigerPlugin
12 - axolotl.integrations.cut_cross_entropy.CutCrossEntropyPlugin
13liger_rope: true
14liger_rms_norm: true
15liger_layer_norm: true
16liger_glu_activation: true
17liger_fused_linear_cross_entropy: false
18cut_cross_entropy: true
19
20load_in_8bit: false
21load_in_4bit: false
22strict: false
23
24datasets:
25 - path: NewEden/RP-logs-V2-Experimental-prefixed
26 type: dan-chat-advanced
27 - path: NewEden/Creative_Writing-Complexity
28 type: dan-chat-advanced
29 - path: NewEden/Discord-Filtered
30 type: dan-chat-advanced
31 - path: NewEden/DeepseekRP-Filtered
32 type: dan-chat-advanced
33 - path: NewEden/Storium-Prefixed-Clean
34 type: dan-chat-advanced
35 - path: NewEden/Basket-Weaving-Filtered
36 type: dan-chat-advanced
37 - path: NewEden/LIMARP-Complexity
38 type: dan-chat-advanced
39 - path: NewEden/Misc-Data-Sharegpt-Prefixed
40 type: dan-chat-advanced
41 - path: NewEden/BlueSky-10K-Complexity
42 type: dan-chat-advanced
43 - path: NewEden/OpenCAI-ShareGPT
44 type: dan-chat-advanced
45 - path: NewEden/Basket-Weaving-Filtered
46 type: dan-chat-advanced
47 - path: PocketDoc/Dans-Personamaxx-VN
48 type: dan-chat-advanced
49 - path: PocketDoc/Dans-Kinomaxx-VanillaBackrooms
50 type: dan-chat-advanced
51dataset_prepared_path: prepared_data
52val_set_size: 0.0
53output_dir: ./qwq-inst
54
55sequence_len: 32768
56sample_packing: true
57pad_to_sequence_len: true
58
59# adapter: lora
60# lora_model_dir:
61# lora_r: 128
62# lora_alpha: 16
63# lora_dropout: 0.05
64# lora_target_modules:
65# - gate_proj
66# - down_proj
67# - up_proj
68# - q_proj
69# - v_proj
70# - k_proj
71# - o_proj
72
73wandb_project: qwq
74wandb_entity:
75wandb_watch:
76wandb_name: rp-attempt-03
77wandb_log_model:
78
79gradient_accumulation_steps: 2
80micro_batch_size: 2
81num_epochs: 4
82optimizer: adamw_bnb_8bit
83lr_scheduler: cosine
84learning_rate: 2.5e-5
85max_grad_norm: 1.0
86
87train_on_inputs: false
88group_by_length: false
89bf16: auto
90fp16:
91tf32: false
92
93gradient_checkpointing: unsloth
94early_stopping_patience:
95resume_from_checkpoint:
96local_rank:
97logging_steps: 1
98xformers_attention:
99flash_attention: true
100
101warmup_steps: 40
102saves_per_epoch: 2
103debug:
104deepspeed: deepspeed_configs/zero3_bf16.json
105weight_decay: 0.02
106fsdp:
107fsdp_config:
108special_tokens: