Views
No views yet

pip install --upgrade autoawq autoawq-kernels1from awq import AutoAWQForCausalLM
2from transformers import AutoTokenizer, TextStreamer
3
4model_path = "solidrust/WestLake-7B-v2-laser-AWQ"
5system_message = "Welcome to WestLake. You are here to help users with any questions they may have."
6
7# Load model
8model = AutoAWQForCausalLM.from_quantized(model_path,
9 fuse_layers=True)
10tokenizer = AutoTokenizer.from_pretrained(model_path,
11 trust_remote_code=True)
12streamer = TextStreamer(tokenizer,
13 skip_prompt=True,
14 skip_special_tokens=True)
15
16# Convert prompt to tokens
17prompt_template = """\
18<|system|>
19</s>
20<|user|>
21{prompt}</s>
22<|assistant|>"""
23
24prompt = "You're standing on the surface of the Earth. "\
25 "You walk one mile south, one mile west and one mile north. "\
26 "You end up exactly where you started. Where are you?"
27
28tokens = tokenizer(prompt_template.format(prompt=prompt),
29 return_tensors='pt').input_ids.cuda()
30
31# Generate output
32generation_output = model.generate(tokens,
33 streamer=streamer,
34 max_new_tokens=512)1<|im_start|>system
2{system_message}<|im_end|>
3<|im_start|>user
4{prompt}<|im_end|>
5<|im_start|>assistant1<|system|>
2</s>
3<|user|>
4{prompt}</s>
5<|assistant|>