|
|
@ -71,6 +71,8 @@ bool gpt_params_parse(int argc, char ** argv, gpt_params & params) {
|
|
|
|
params.use_color = true;
|
|
|
|
params.use_color = true;
|
|
|
|
} else if (arg == "-r" || arg == "--reverse-prompt") {
|
|
|
|
} else if (arg == "-r" || arg == "--reverse-prompt") {
|
|
|
|
params.antiprompt = argv[++i];
|
|
|
|
params.antiprompt = argv[++i];
|
|
|
|
|
|
|
|
} else if (arg == "--ignore-eos") {
|
|
|
|
|
|
|
|
params.ignore_eos = true;
|
|
|
|
} else if (arg == "-h" || arg == "--help") {
|
|
|
|
} else if (arg == "-h" || arg == "--help") {
|
|
|
|
gpt_print_usage(argc, argv, params);
|
|
|
|
gpt_print_usage(argc, argv, params);
|
|
|
|
exit(0);
|
|
|
|
exit(0);
|
|
|
@ -106,6 +108,7 @@ void gpt_print_usage(int /*argc*/, char ** argv, const gpt_params & params) {
|
|
|
|
fprintf(stderr, " --repeat_last_n N last n tokens to consider for penalize (default: %d)\n", params.repeat_last_n);
|
|
|
|
fprintf(stderr, " --repeat_last_n N last n tokens to consider for penalize (default: %d)\n", params.repeat_last_n);
|
|
|
|
fprintf(stderr, " --repeat_penalty N penalize repeat sequence of tokens (default: %.1f)\n", params.repeat_penalty);
|
|
|
|
fprintf(stderr, " --repeat_penalty N penalize repeat sequence of tokens (default: %.1f)\n", params.repeat_penalty);
|
|
|
|
fprintf(stderr, " -c N, --ctx_size N size of the prompt context (default: %d)\n", params.n_ctx);
|
|
|
|
fprintf(stderr, " -c N, --ctx_size N size of the prompt context (default: %d)\n", params.n_ctx);
|
|
|
|
|
|
|
|
fprintf(stderr, " --ignore-eos ignore end of stream token and continue generating\n");
|
|
|
|
fprintf(stderr, " --memory_f16 use f16 instead of f32 for memory key+value\n");
|
|
|
|
fprintf(stderr, " --memory_f16 use f16 instead of f32 for memory key+value\n");
|
|
|
|
fprintf(stderr, " --temp N temperature (default: %.1f)\n", params.temp);
|
|
|
|
fprintf(stderr, " --temp N temperature (default: %.1f)\n", params.temp);
|
|
|
|
fprintf(stderr, " -b N, --batch_size N batch size for prompt processing (default: %d)\n", params.n_batch);
|
|
|
|
fprintf(stderr, " -b N, --batch_size N batch size for prompt processing (default: %d)\n", params.n_batch);
|
|
|
|