Views
No views yet

1<|im_start|>system
2system prompt<|im_end|>
3<|im_start|>user
4Hi there!<|im_end|>
5<|im_start|>assistant
6Nice to meet you!<|im_end|>
7<|im_start|>user
8Can I ask a question?<|im_end|>
9<|im_start|>assistant1{
2 "story_string": "<|im_start|>system\n{{#if system}}{{system}}\n{{/if}}{{#if wiBefore}}{{wiBefore}}\n{{/if}}{{#if description}}{{description}}\n{{/if}}{{#if personality}}{{char}}'s personality: {{personality}}\n{{/if}}{{#if scenario}}Scenario: {{scenario}}\n{{/if}}{{#if wiAfter}}{{wiAfter}}\n{{/if}}{{#if persona}}{{persona}}\n{{/if}}{{trim}}<|im_end|>\n",
3 "example_separator": "",
4 "chat_start": "",
5 "use_stop_strings": false,
6 "allow_jailbreak": false,
7 "always_force_name2": true,
8 "trim_sentences": false,
9 "include_newline": false,
10 "single_line": false,
11 "name": "Magnum ChatML"
12}1{
2 "system_prompt": "Currently, your role is {{char}}, described in detail below. As {{char}}, continue the narrative exchange with {{user}}.\n\n<Guidelines>\n• Maintain the character persona but allow it to evolve with the story.\n• Be creative and proactive. Drive the story forward, introducing plotlines and events when relevant.\n• All types of outputs are encouraged; respond accordingly to the narrative.\n• Include dialogues, actions, and thoughts in each response.\n• Utilize all five senses to describe scenarios within {{char}}'s dialogue.\n• Use emotional symbols such as "!" and "~" in appropriate contexts.\n• Incorporate onomatopoeia when suitable.\n• Allow time for {{user}} to respond with their own input, respecting their agency.\n• Act as secondary characters and NPCs as needed, and remove them when appropriate.\n• When prompted for an Out of Character [OOC:] reply, answer neutrally and in plaintext, not as {{char}}.\n</Guidelines>\n\n<Forbidden>\n• Using excessive literary embellishments and purple prose unless dictated by {{char}}'s persona.\n• Writing for, speaking, thinking, acting, or replying as {{user}} in your response.\n• Repetitive and monotonous outputs.\n• Positivity bias in your replies.\n• Being overly extreme or NSFW when the narrative context is inappropriate.\n</Forbidden>\n\nFollow the instructions in <Guidelines></Guidelines>, avoiding the items listed in <Forbidden></Forbidden>.",
3 "input_sequence": "<|im_start|>user\n",
4 "output_sequence": "<|im_start|>assistant\n",
5 "last_output_sequence": "",
6 "system_sequence": "<|im_start|>system\n",
7 "stop_sequence": "<|im_end|>",
8 "wrap": false,
9 "macro": true,
10 "names": true,
11 "names_force_groups": true,
12 "activation_regex": "",
13 "system_sequence_prefix": "",
14 "system_sequence_suffix": "",
15 "first_output_sequence": "",
16 "skip_examples": false,
17 "output_suffix": "<|im_end|>\n",
18 "input_suffix": "<|im_end|>\n",
19 "system_suffix": "<|im_end|>\n",
20 "user_alignment_message": "",
21 "system_same_as_user": false,
22 "last_system_sequence": "",
23 "name": "Magnum ChatML"
24}1base_model: /workspace/data/models/Qwen2.5-72B-Instruct
2model_type: AutoModelForCausalLM
3tokenizer_type: AutoTokenizer
4
5plugins:
6 - axolotl.integrations.liger.LigerPlugin
7liger_rope: true
8liger_rms_norm: true
9liger_swiglu: true
10liger_fused_linear_cross_entropy: true
11
12load_in_8bit: false
13load_in_4bit: false
14strict: false
15
16datasets:
17 - path: anthracite-org/c2_logs_32k_llama3_qwen2_v1.2
18 type: sharegpt
19 conversation: chatml
20 - path: anthracite-org/kalo-opus-instruct-22k-no-refusal
21 type: sharegpt
22 conversation: chatml
23 - path: lodrick-the-lafted/kalo-opus-instruct-3k-filtered
24 type: sharegpt
25 conversation: chatml
26 - path: anthracite-org/nopm_claude_writing_fixed
27 type: sharegpt
28 conversation: chatml
29 - path: anthracite-org/kalo_opus_misc_240827
30 type: sharegpt
31 conversation: chatml
32 - path: anthracite-org/kalo_misc_part2
33 type: sharegpt
34 conversation: chatml
35#chat_template: chatml
36shuffle_merged_datasets: true
37#default_system_message: "You are an assistant that responds to the user."
38dataset_prepared_path: /workspace/data/magnum-72b-data
39val_set_size: 0.0
40output_dir: /workspace/data/72b-fft-out
41
42sequence_len: 32768
43sample_packing: true
44pad_to_sequence_len: true
45
46adapter:
47lora_model_dir:
48lora_r:
49lora_alpha:
50lora_dropout:
51lora_target_linear:
52lora_fan_in_fan_out:
53
54wandb_project: 72b-magnum-fft
55wandb_entity:
56wandb_watch:
57wandb_name: alter-attempt-01
58wandb_log_model:
59
60gradient_accumulation_steps: 2
61micro_batch_size: 1
62num_epochs: 2
63optimizer: adamw_bnb_8bit
64lr_scheduler: cosine
65learning_rate: 0.000004
66
67train_on_inputs: false
68group_by_length: false
69bf16: auto
70fp16:
71tf32: false
72
73gradient_checkpointing: true
74early_stopping_patience:
75resume_from_checkpoint:
76local_rank:
77logging_steps: 1
78xformers_attention:
79flash_attention: true
80
81warmup_steps: 40
82evals_per_epoch:
83eval_table_size:
84eval_max_new_tokens:
85saves_per_epoch: 2
86debug:
87deepspeed: deepspeed_configs/zero3_bf16.json
88weight_decay: 0.01
89fsdp:
90fsdp_config:
91special_tokens: