import torch
from transformers import AutoTokenizer, AutoModelForCausalLM
tokenizer = AutoTokenizer.from_pretrained("rinna/nekomata-7b-instruction", trust_remote_code=True)
# Use GPU with bf16# model = AutoModelForCausalLM.from_pretrained("rinna/nekomata-7b-instruction", device_map="auto", trust_remote_code=True, bf16=True)# Use GPU with fp16# model = AutoModelForCausalLM.from_pretrained("rinna/nekomata-7b-instruction", device_map="auto", trust_remote_code=True, fp16=True)# Use CPU# model = AutoModelForCausalLM.from_pretrained("rinna/nekomata-7b-instruction", device_map="cpu", trust_remote_code=True)# Automatically select device and precision
model = AutoModelForCausalLM.from_pretrained("rinna/nekomata-7b-instruction", device_map="auto", trust_remote_code=True)
instruction = "次の日本語を英語に翻訳してください。"input = "大規模言語モデル(だいきぼげんごモデル、英: large language model、LLM)は、多数のパラメータ(数千万から数十億)を持つ人工ニューラルネットワークで構成されるコンピュータ言語モデルで、膨大なラベルなしテキストを使用して自己教師あり学習または半教師あり学習によって訓練が行われる。"
prompt = f"""以下は、タスクを説明する指示と、文脈のある入力の組み合わせです。要求を適切に満たす応答を書きなさい。### 指示:{instruction}### 入力:{input}### 応答:"""
token_ids = tokenizer.encode(prompt, add_special_tokens=False, return_tensors="pt")
with torch.no_grad():
output_ids = model.generate(
token_ids.to(model.device),
max_new_tokens=200,
do_sample=True,
temperature=0.5,
pad_token_id=tokenizer.pad_token_id,
bos_token_id=tokenizer.bos_token_id,
eos_token_id=tokenizer.eos_token_id
)
output = tokenizer.decode(output_ids.tolist()[0])
print(output)
"""以下は、タスクを説明する指示と、文脈のある入力の組み合わせです。要求を適切に満たす応答を書きなさい。### 指示:次の日本語を英語に翻訳してください。### 入力:大規模言語モデル(だいきぼげんごモデル、英: large language model、LLM)は、多数のパラメータ(数千万から数十億)を持つ人工ニューラルネットワークで構成されるコンピュータ言語モデルで、膨大なラベルなしテキストを使 用して自己教師あり学習または半教師あり学習によって訓練が行われる。### 応答: A large language model (LLM) is a computer language model composed of artificial neural networks with many parameters (from tens of millions to billions) trained by self-supervised learning or semi-supervised learning using a large amount of unlabeled text.<|endoftext|>"""
@misc{rinna-nekomata-7b-instruction,
title = {rinna/nekomata-7b-instruction},
author = {Zhao, Tianyu and Sawada, Kei},
url = {https://huggingface.co/rinna/nekomata-7b-instruction}
}
@inproceedings{sawada2024release,
title = {Release of Pre-Trained Models for the {J}apanese Language},
author = {Sawada, Kei and Zhao, Tianyu and Shing, Makoto and Mitsui, Kentaro and Kaga, Akio and Hono, Yukiya and Wakatsuki, Toshiaki and Mitsuda, Koh},
booktitle = {Proceedings of the 2024 Joint International Conference on Computational Linguistics, Language Resources and Evaluation (LREC-COLING 2024)},
month = {5},
year = {2024},
pages = {13898--13905},
url = {https://aclanthology.org/2024.lrec-main.1213},
note = {\url{https://arxiv.org/abs/2404.01657}}
}
nekomata-7b-instruction huggingface.co is an AI model on huggingface.co that provides nekomata-7b-instruction's model effect (), which can be used instantly with this rinna nekomata-7b-instruction model. huggingface.co supports a free trial of the nekomata-7b-instruction model, and also provides paid use of the nekomata-7b-instruction. Support call nekomata-7b-instruction model through api, including Node.js, Python, http.
nekomata-7b-instruction huggingface.co is an online trial and call api platform, which integrates nekomata-7b-instruction's modeling effects, including api services, and provides a free online trial of nekomata-7b-instruction, you can try nekomata-7b-instruction online for free by clicking the link below.
rinna nekomata-7b-instruction online free url in huggingface.co:
nekomata-7b-instruction is an open source model from GitHub that offers a free installation service, and any user can find nekomata-7b-instruction on GitHub to install. At the same time, huggingface.co provides the effect of nekomata-7b-instruction install, users can directly use nekomata-7b-instruction installed effect in huggingface.co for debugging and trial. It also supports api for free installation.
nekomata-7b-instruction install url in huggingface.co: