mirror of
https://github.com/deepseek-ai/DeepSeek-V3
synced 2025-01-22 12:25:30 +00:00
require model-parallel in convert.py
This commit is contained in:
parent
7c2466b310
commit
8710ec2ecb
@ -78,7 +78,7 @@ if __name__ == "__main__":
|
|||||||
parser.add_argument("--hf-ckpt-path", type=str, required=True)
|
parser.add_argument("--hf-ckpt-path", type=str, required=True)
|
||||||
parser.add_argument("--save-path", type=str, required=True)
|
parser.add_argument("--save-path", type=str, required=True)
|
||||||
parser.add_argument("--n-experts", type=int, required=True)
|
parser.add_argument("--n-experts", type=int, required=True)
|
||||||
parser.add_argument("--model-parallel", type=int, default=1)
|
parser.add_argument("--model-parallel", type=int, required=True)
|
||||||
args = parser.parse_args()
|
args = parser.parse_args()
|
||||||
assert args.n_experts % args.model_parallel == 0
|
assert args.n_experts % args.model_parallel == 0
|
||||||
main(args.hf_ckpt_path, args.save_path, args.n_experts, args.model_parallel)
|
main(args.hf_ckpt_path, args.save_path, args.n_experts, args.model_parallel)
|
||||||
|
Loading…
Reference in New Issue
Block a user