Spaces:

mikeee
/

chatglm2-6b-test

Runtime error

App Files Files Community

ffreemt commited on Jul 14, 2023

Commit

634ed9b

1 Parent(s): d82d66e

Update layout

Browse files

Files changed (1) hide show

app.py +58 -10

app.py CHANGED Viewed

@@ -15,7 +15,6 @@ from transformers import AutoModel, AutoTokenizer
 # os.system("pip install torch transformers sentencepiece loguru")
 # fix timezone in Linux
 os.environ["TZ"] = "Asia/Shanghai"
 try:
@@ -42,9 +41,9 @@ if has_cuda:
             AutoModel.from_pretrained(model_name, trust_remote_code=True).cuda().half()
         )
 else:
-    model = AutoModel.from_pretrained(
-        model_name, trust_remote_code=True
-    ).half().float()  # .float() .half().float(): must use float for cpu
 model = model.eval()
 logger.debug("done load")
@@ -54,7 +53,11 @@ logger.debug("done load")
 # locate model file cache
 cache_loc = Path("~/.cache/huggingface/hub").expanduser()
-model_cache_path = [elm for elm in Path(cache_loc).rglob("*") if Path(model_name).name in elm.as_posix() and "pytorch_model.bin" in elm.as_posix()]
 logger.debug(f"{model_cache_path=}")
@@ -64,25 +67,70 @@ if model_cache_path:
 def respond(message, chat_history):
-    response, chat_history = model.chat(tokenizer, message, history=chat_history, temperature=0.7, repetition_penalty=1.2, max_length=128)
     chat_history.append((message, response))
-    return "", chat_history
 theme = gr.themes.Soft(text_size="sm")
 with gr.Blocks(theme=theme) as block:
     chatbot = gr.Chatbot()
-    with gr.Column():
         with gr.Column(scale=12):
             msg = gr.Textbox()
         with gr.Column(scale=1, min_width=16):
-            with gr.Row():
                 btn = gr.Button("Send")
                 clear = gr.ClearButton([msg, chatbot])
     # do not clear prompt
-    msg.submit(lambda x, y: [x] + respond(x, y)[1:], [msg, chatbot], [msg, chatbot])
     btn.click(respond, [msg, chatbot], [msg, chatbot])
 block.queue().launch()

 # os.system("pip install torch transformers sentencepiece loguru")
 # fix timezone in Linux
 os.environ["TZ"] = "Asia/Shanghai"
 try:
             AutoModel.from_pretrained(model_name, trust_remote_code=True).cuda().half()
         )
 else:
+    model = (
+        AutoModel.from_pretrained(model_name, trust_remote_code=True).half().float()
+    )  # .float() .half().float(): must use float for cpu
 model = model.eval()
 logger.debug("done load")
 # locate model file cache
 cache_loc = Path("~/.cache/huggingface/hub").expanduser()
+model_cache_path = [
+    elm
+    for elm in Path(cache_loc).rglob("*")
+    if Path(model_name).name in elm.as_posix() and "pytorch_model.bin" in elm.as_posix()
+]
 logger.debug(f"{model_cache_path=}")
 def respond(message, chat_history):
+    """Gen a response."""
+    message = message.strip()
+    response, chat_history = model.chat(
+        tokenizer,
+        message,
+        history=chat_history,
+        temperature=0.7,
+        repetition_penalty=1.2,
+        max_length=128,
+    )
     chat_history.append((message, response))
+    return message, chat_history
 theme = gr.themes.Soft(text_size="sm")
 with gr.Blocks(theme=theme) as block:
     chatbot = gr.Chatbot()
+    with gr.Row():
         with gr.Column(scale=12):
             msg = gr.Textbox()
         with gr.Column(scale=1, min_width=16):
                 btn = gr.Button("Send")
+        with gr.Column(scale=1, min_width=8):
                 clear = gr.ClearButton([msg, chatbot])
     # do not clear prompt
+    msg.submit(lambda x, y: (x,) + respond(x, y)[1:], [msg, chatbot], [msg, chatbot])
     btn.click(respond, [msg, chatbot], [msg, chatbot])
+    with gr.Accordion("Example inputs", open=True):
+        etext = """In America, where cars are an important part of the national psyche, a decade ago people had suddenly started to drive less, which had not happened since the oil shocks of the 1970s. """
+        examples = gr.Examples(
+            examples=[
+                ["Explain the plot of Cinderella in a sentence."],
+                [
+                    "How long does it take to become proficient in French, and what are the best methods for retaining information?"
+                ],
+                ["What are some common mistakes to avoid when writing code?"],
+                ["Build a prompt to generate a beautiful portrait of a horse"],
+                ["Suggest four metaphors to describe the benefits of AI"],
+                ["Write a pop song about leaving home for the sandy beaches."],
+                ["Write a summary demonstrating my ability to tame lions"],
+                ["鲁迅和周树人什么关系"],
+                ["从前有一头牛，这头牛后面有什么？"],
+                ["正无穷大加一大于正无穷大吗？"],
+                ["正无穷大加正无穷大大于正无穷大吗？"],
+                ["-2的平方根等于什么"],
+                ["树上有5只鸟，猎人开枪打死了一只。树上还有几只鸟？"],
+                ["树上有11只鸟，猎人开枪打死了一只。树上还有几只鸟？提示：需考虑鸟可能受惊吓飞走。"],
+                ["鲁迅和周树人什么关系 用英文回答"],
+                ["以红楼梦的行文风格写一张委婉的请假条。不少于320字。"],
+                [f"{etext} 翻成中文，列出3个版本"],
+                [f"{etext} \n 翻成中文，保留原意，但使用文学性的语言。不要写解释。列出3个版本"],
+                ["js 判断一个数是不是质数"],
+                ["js 实现python 的 range(10)"],
+                ["js 实现python 的 [*(range(10)]"],
+                ["假定 1 + 2 = 4, 试求 7 + 8"],
+                ["Erkläre die Handlung von Cinderella in einem Satz."],
+                ["Erkläre die Handlung von Cinderella in einem Satz. Auf Deutsch"],
+            ],
+            inputs=[msg],
+            examples_per_page=60,
+        )
 block.queue().launch()