Eclipse-Senpai commited on
Commit
495e2c4
·
verified ·
1 Parent(s): e25b8ed

scrub internal project references from generate.py

Browse files
Files changed (1) hide show
  1. generate.py +2 -2
generate.py CHANGED
@@ -1,6 +1,6 @@
1
  """Raw bundled inference CLI for min-spark (no transformers dependency).
2
 
3
- Mirrors spaces/min-spark-preview/loader.py's generation loop exactly: EOS
4
  prefix once, truncate to the last max_seq_len tokens, effort -> loop count.
5
  This is the second, self-contained integration path; the Transformers path is
6
  modeling_minspark.py. Prefer the Transformers path unless you want zero
@@ -36,7 +36,7 @@ def load_model(ckpt_path: str | None = None, device: str = "cpu") -> Meiosis:
36
  @torch.no_grad()
37
  def generate(model, tokenizer, prompt: str, *, loops: int, max_new: int,
38
  temperature: float, top_k: int, device: str):
39
- """Yield decoded tokens one at a time (mirrors loader.py)."""
40
  ids = [EOS_ID] + tokenizer.encode(prompt).ids
41
  for _ in range(max_new):
42
  ctx = ids[-model.config.max_seq_len:]
 
1
  """Raw bundled inference CLI for min-spark (no transformers dependency).
2
 
3
+ Mirrors the Space loader's generation loop exactly: EOS
4
  prefix once, truncate to the last max_seq_len tokens, effort -> loop count.
5
  This is the second, self-contained integration path; the Transformers path is
6
  modeling_minspark.py. Prefer the Transformers path unless you want zero
 
36
  @torch.no_grad()
37
  def generate(model, tokenizer, prompt: str, *, loops: int, max_new: int,
38
  temperature: float, top_k: int, device: str):
39
+ """Yield decoded tokens one at a time (mirrors the Space loader)."""
40
  ids = [EOS_ID] + tokenizer.encode(prompt).ids
41
  for _ in range(max_new):
42
  ctx = ids[-model.config.max_seq_len:]