CoquiTTS

Sleeping

App Files Files Community

Eren Gölge commited on May 18, 2023

Commit

a3635dc

1 Parent(s): 8f9ad5e

Update UI

Browse files

Files changed (1) hide show

app.py +114 -62

app.py CHANGED Viewed

@@ -28,7 +28,7 @@ MODEL_NAMES = EN + OTHER
 print(MODEL_NAMES)
-def tts(text: str, model_name: str, speaker_idx: str=None):
     if len(text) > MAX_TXT_LEN:
         text = text[:MAX_TXT_LEN]
         print(f"Input text was cutoff since it went over the {MAX_TXT_LEN} character limit.")
@@ -48,72 +48,124 @@ def tts(text: str, model_name: str, speaker_idx: str=None):
     # synthesize
     if synthesizer is None:
         raise NameError("model not found")
-    wavs = synthesizer.tts(text, speaker_idx)
     # return output
     with tempfile.NamedTemporaryFile(suffix=".wav", delete=False) as fp:
         synthesizer.save_wav(wavs, fp)
         return fp.name
-article= """
-Visit us on Coqui.ai and drop a 🌟 to 🔗<a href="https://github.com/coqui-ai/TTS" target="_blank">CoquiTTS</a>.
-<br/>
-Run CoquiTTS locally for the best result. Check out our 🔗<a href="https://tts.readthedocs.io/en/latest/inference.html">documentation</a>.
-```bash
-$ pip install TTS
-...
-$ tts --list_models
-...
-$ tts --text "Text for TTS" --model_name "<type>/<language>/<dataset>/<model_name>" --out_path folder/to/save/output.wav
-```
-<img src="https://static.scarf.sh/a.png?x-pxid=1404a024-e647-4406-bb9a-4ade0c931182" />
-<br/>
-👑 <b> Model contributors</b>
-- <a href="https://github.com/nmstoker/" target="_blank">@nmstoker</a>
-- <a href="https://github.com/kaiidams/" target="_blank">@kaiidams</a>
-- <a href="https://github.com/WeberJulian/" target="_blank">@WeberJulian,</a>
-- <a href="https://github.com/Edresson/" target="_blank">@Edresson</a>
-- <a href="https://github.com/thorstenMueller/" target="_blank">@thorstenMueller</a>
-- <a href="https://github.com/r-dh/" target="_blank">@r-dh</a>
-- <a href="https://github.com/kirianguiller/" target="_blank">@kirianguiller</a>
-- <a href="https://github.com/robinhad/" target="_blank">@robinhad</a>
-- <a href="https://github.com/fkarabiber/" target="_blank">@fkarabiber</a>
-- <a href="https://github.com/nicolalandro/" target="_blank">@nicolalandro</a>
-- <a href="https://github.com/a-froghyar" target="_blank">@a-froghyar</a>
-- <a href="https://github.com/manmay-nakhashi" target="_blank">@manmay-nakhashi</a>
-- <a href="https://github.com/noml4u" target="_blank">@noml4u</a>
-👉 Drop a ✨PR✨ on 🐸TTS to share a new model and have it included here.
 """
-iface = gr.Interface(
-    fn=tts,
-    inputs=[
-        gr.inputs.Textbox(
-            label="Input Text",
-            default="This sentence has been generated by a speech synthesis system.",
-        ),
-        gr.inputs.Radio(
-            label="Pick a TTS Model - (language/dataset/model_name)",
-            choices=MODEL_NAMES,
-        ),
-        # gr.inputs.Dropdown(label="Select a speaker", choices=SPEAKERS, default=None)
-        # gr.inputs.Audio(source="microphone", label="Record your voice.", type="numpy", label=None, optional=False)
-    ],
-    outputs=gr.outputs.Audio(label="Output"),
-    title="🐸💬 CoquiTTS Demo",
-    theme="grass",
-    description="🐸💬  Coqui TTS - a deep learning toolkit for Text-to-Speech, battle-tested in research and production.",
-    article=article,
-    allow_flagging=False,
-    flagging_options=['error', 'bad-quality', 'wrong-pronounciation'],
-    layout="vertical",
-    live=False
-)
-iface.launch(share=False)

 print(MODEL_NAMES)
+def tts(text: str, model_name: str):
     if len(text) > MAX_TXT_LEN:
         text = text[:MAX_TXT_LEN]
         print(f"Input text was cutoff since it went over the {MAX_TXT_LEN} character limit.")
     # synthesize
     if synthesizer is None:
         raise NameError("model not found")
+    wavs = synthesizer.tts(text, None)
     # return output
     with tempfile.NamedTemporaryFile(suffix=".wav", delete=False) as fp:
         synthesizer.save_wav(wavs, fp)
         return fp.name
+title = """<h1 align="center">🐸💬 CoquiTTS Playground </h1>"""
+custom_css = """
+#banner-image {
+    display: block;
+    margin-left: auto;
+    margin-right: auto;
+}
 """
+with gr.Blocks(analytics_enabled=False, css=custom_css) as demo:
+    with gr.Row():
+        with gr.Column():
+            # gr.Image("https://raw.githubusercontent.com/coqui-ai/TTS/main/images/coqui-log-green-TTS.png", elem_id="banner-image", show_label=False)
+            gr.Markdown(
+                """
+                ## <img src="https://raw.githubusercontent.com/coqui-ai/TTS/main/images/coqui-log-green-TTS.png" height="56"/>
+                """
+            )
+            gr.Markdown(
+                """
+                <br />
+                ## 🐸Coqui.ai News
+                - 📣 🐸TTS now supports 🐢Tortoise with faster inference.
+                - 📣 **Coqui Studio API** is landed on 🐸TTS. - [Example](https://github.com/coqui-ai/TTS/blob/dev/README.md#-python-api)
+                - 📣 [**Coqui Sudio API**](https://docs.coqui.ai/docs) is live.
+                - 📣 Voice generation with prompts - **Prompt to Voice** - is live on [**Coqui Studio**](https://app.coqui.ai/auth/signin)!! - [Blog Post](https://coqui.ai/blog/tts/prompt-to-voice)
+                - 📣 Voice generation with fusion - **Voice fusion** - is live on [**Coqui Studio**](https://app.coqui.ai/auth/signin).
+                - 📣 Voice cloning is live on [**Coqui Studio**](https://app.coqui.ai/auth/signin).
+                <br>
+                """
+                )
+        with gr.Column():
+            gr.Markdown(
+            """
+            <br/>
+            💻 This space showcases some of the **[CoquiTTS](https://github.com/coqui-ai/TTS)** models.
+            <br/>
+            There are > 30 languages with single and multi speaker models, all thanks to our 👑 Contributors.
+            <br/>
+            Visit the links below for more.
+            |                                 |                                         |
+            | ------------------------------- | --------------------------------------- |
+            | 🐸💬 **CoquiTTS**                |  [Github](https://github.com/coqui-ai/TTS) |
+            | 💼 **Documentation**            | [ReadTheDocs](https://tts.readthedocs.io/en/latest/)
+            | 👩‍💻 **Questions**                | [GitHub Discussions] |
+            | 🗯 **Community**         | [![Dicord](https://img.shields.io/discord/1037326658807533628?color=%239B59B6&label=chat%20on%20discord)](https://discord.gg/5eXr5seRrv)  |
+            [github issue tracker]: https://github.com/coqui-ai/tts/issues
+            [github discussions]: https://github.com/coqui-ai/TTS/discussions
+            [discord]: https://discord.gg/5eXr5seRrv
+            """
+            )
+    with gr.Row():
+        gr.Markdown(
+            """
+            <details>
+            <summary>👑 Model contributors</summary>
+            - <a href="https://github.com/nmstoker/" target="_blank">@nmstoker</a>
+            - <a href="https://github.com/kaiidams/" target="_blank">@kaiidams</a>
+            - <a href="https://github.com/WeberJulian/" target="_blank">@WeberJulian,</a>
+            - <a href="https://github.com/Edresson/" target="_blank">@Edresson</a>
+            - <a href="https://github.com/thorstenMueller/" target="_blank">@thorstenMueller</a>
+            - <a href="https://github.com/r-dh/" target="_blank">@r-dh</a>
+            - <a href="https://github.com/kirianguiller/" target="_blank">@kirianguiller</a>
+            - <a href="https://github.com/robinhad/" target="_blank">@robinhad</a>
+            - <a href="https://github.com/fkarabiber/" target="_blank">@fkarabiber</a>
+            - <a href="https://github.com/nicolalandro/" target="_blank">@nicolalandro</a>
+            - <a href="https://github.com/a-froghyar" target="_blank">@a-froghyar</a>
+            - <a href="https://github.com/manmay-nakhashi" target="_blank">@manmay-nakhashi</a>
+            - <a href="https://github.com/noml4u" target="_blank">@noml4u</a>
+            </details>
+            <br/>
+            """
+        )
+    with gr.Row():
+        with gr.Column():
+            input_text = gr.inputs.Textbox(
+                label="Input Text",
+                default="This sentence has been generated by a speech synthesis system.",
+            )
+            model_select = gr.inputs.Dropdown(
+                label="Pick Model: tts_models/<language>/<dataset>/<model_name>",
+                choices=MODEL_NAMES,
+                default="tts_models/en/jenny/jenny"
+            )
+            tts_button = gr.Button("Send", elem_id="send-btn", visible=True)
+        with gr.Column():
+            output_audio = gr.outputs.Audio(label="Output", type="filepath")
+    tts_button.click(
+        tts,
+        inputs=[
+            input_text,
+            model_select,
+        ],
+        outputs=[output_audio],
+    )
+demo.queue(concurrency_count=16).launch(debug=True)