Main+: optionally allow special tokens from user in interactive mode (#7097)

author HanishKVC <redacted>

Fri, 10 May 2024 10:21:58 +0000 (15:51 +0530)

committer GitHub <redacted>

Fri, 10 May 2024 10:21:58 +0000 (20:21 +1000)
author HanishKVC <redacted>
Fri, 10 May 2024 10:21:58 +0000 (15:51 +0530)
committer GitHub <redacted>
Fri, 10 May 2024 10:21:58 +0000 (20:21 +1000)
diff --git a/common/common.cpp b/common/common.cpp

index 0535508ba98dfafd8ce5658b278160b20fb825a4..484e673349071b77b8bbdbd5b1d1a45513a82ba1 100644 (file)
--- a/common/common.cpp
+++ b/common/common.cpp
@@ -901,6 +901,10 @@ bool gpt_params_find_arg(int argc, char ** argv, const std::string & arg, gpt_pa
          params.interactive = true;
          return true;
      }
+    if (arg == "--interactive-specials") {
+        params.interactive_specials = true;
+        return true;
+    }
      if (arg == "--embedding") {
          params.embedding = true;
          return true;
@@ -1422,6 +1426,7 @@ void gpt_print_usage(int /*argc*/, char ** argv, const gpt_params & params) {
      printf("  -h, --help            show this help message and exit\n");
      printf("  --version             show version and build info\n");
      printf("  -i, --interactive     run in interactive mode\n");
+    printf("  --interactive-specials allow special tokens in user text, in interactive mode\n");
      printf("  --interactive-first   run in interactive mode and wait for input right away\n");
      printf("  -cnv, --conversation  run in conversation mode (does not print special tokens and suffix/prefix)\n");
      printf("  -ins, --instruct      run in instruction mode (use with Alpaca models)\n");
@@ -2652,6 +2657,7 @@ void dump_non_result_info_yaml(FILE * stream, const gpt_params & params, const l
      dump_string_yaml_multiline(stream, "in_suffix", params.input_prefix.c_str());
      fprintf(stream, "instruct: %s # default: false\n", params.instruct ? "true" : "false");
      fprintf(stream, "interactive: %s # default: false\n", params.interactive ? "true" : "false");
+    fprintf(stream, "interactive_specials: %s # default: false\n", params.interactive_specials ? "true" : "false");
      fprintf(stream, "interactive_first: %s # default: false\n", params.interactive_first ? "true" : "false");
      fprintf(stream, "keep: %d # default: 0\n", params.n_keep);
      fprintf(stream, "logdir: %s # default: unset (no logging)\n", params.logdir.c_str());
diff --git a/common/common.h b/common/common.h

index 6f00a2cca888374093224a5163339ed2d35ccd2c..d80344f2a61bbc39bc0247f840455fd100b15eec 100644 (file)
--- a/common/common.h
+++ b/common/common.h
@@ -140,6 +140,7 @@ struct gpt_params {
      bool random_prompt     = false; // do not randomize prompt if none provided
      bool use_color         = false; // use color to distinguish generations and inputs
      bool interactive       = false; // interactive mode
+    bool interactive_specials = false; // whether to allow special tokens from user, during interactive mode
      bool conversation      = false; // conversation mode (does not print special tokens and suffix/prefix)
      bool chatml            = false; // chatml mode (used for models trained on chatml syntax)
      bool prompt_cache_all  = false; // save user input and generations to prompt cache
diff --git a/examples/main/main.cpp b/examples/main/main.cpp

index 49acd6bab4074ac5bf23e88fa588f6c114dd17e8..f3e445c16d6a9ba5512f27f00a5a41cc0a38bf1a 100644 (file)
--- a/examples/main/main.cpp
+++ b/examples/main/main.cpp
@@ -879,7 +879,7 @@ int main(int argc, char ** argv) {
                      }
  
                      const auto line_pfx = ::llama_tokenize(ctx, params.input_prefix, false, true);
-                    const auto line_inp = ::llama_tokenize(ctx, buffer,              false, false);
+                    const auto line_inp = ::llama_tokenize(ctx, buffer,              false, params.interactive_specials);
                      const auto line_sfx = ::llama_tokenize(ctx, params.input_suffix, false, true);
  
                      LOG("input tokens: %s\n", LOG_TOKENS_TOSTR_PRETTY(ctx, line_inp).c_str());
author	HanishKVC <redacted>
	Fri, 10 May 2024 10:21:58 +0000 (15:51 +0530)
committer	GitHub <redacted>
	Fri, 10 May 2024 10:21:58 +0000 (20:21 +1000)
common/common.cpp		patch \| blob \| history
common/common.h		patch \| blob \| history
examples/main/main.cpp		patch \| blob \| history