mirror of
https://github.com/ml-explore/mlx-examples.git
synced 2025-08-28 07:54:24 +08:00
fix prompt cache with no chat template
This commit is contained in:
parent
6e6ba07b54
commit
087adcfacb
@ -152,7 +152,7 @@ def main():
|
|||||||
print("Saving...")
|
print("Saving...")
|
||||||
metadata = {}
|
metadata = {}
|
||||||
metadata["model"] = args.model
|
metadata["model"] = args.model
|
||||||
metadata["chat_template"] = tokenizer.chat_template
|
metadata["chat_template"] = json.dumps(tokenizer.chat_template)
|
||||||
metadata["tokenizer_config"] = json.dumps(tokenizer_config)
|
metadata["tokenizer_config"] = json.dumps(tokenizer_config)
|
||||||
save_prompt_cache(args.prompt_cache_file, cache, metadata)
|
save_prompt_cache(args.prompt_cache_file, cache, metadata)
|
||||||
|
|
||||||
|
@ -199,7 +199,7 @@ def main():
|
|||||||
if tokenizer.chat_template is None:
|
if tokenizer.chat_template is None:
|
||||||
tokenizer.chat_template = tokenizer.default_chat_template
|
tokenizer.chat_template = tokenizer.default_chat_template
|
||||||
elif using_cache:
|
elif using_cache:
|
||||||
tokenizer.chat_template = metadata["chat_template"]
|
tokenizer.chat_template = json.loads(metadata["chat_template"])
|
||||||
|
|
||||||
prompt = args.prompt.replace("\\n", "\n").replace("\\t", "\t")
|
prompt = args.prompt.replace("\\n", "\n").replace("\\t", "\t")
|
||||||
prompt = sys.stdin.read() if prompt == "-" else prompt
|
prompt = sys.stdin.read() if prompt == "-" else prompt
|
||||||
|
Loading…
Reference in New Issue
Block a user