2 * mdadm - manage Linux "md" devices aka RAID arrays.
4 * Copyright (C) 2001-2013 Neil Brown <neilb@suse.de>
7 * This program is free software; you can redistribute it and/or modify
8 * it under the terms of the GNU General Public License as published by
9 * the Free Software Foundation; either version 2 of the License, or
10 * (at your option) any later version.
12 * This program is distributed in the hope that it will be useful,
13 * but WITHOUT ANY WARRANTY; without even the implied warranty of
14 * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the
15 * GNU General Public License for more details.
17 * You should have received a copy of the GNU General Public License
18 * along with this program; if not, write to the Free Software
19 * Foundation, Inc., 59 Temple Place, Suite 330, Boston, MA 02111-1307 USA
22 * Email: <neilb@suse.de>
24 * Additions for bitmap and write-behind RAID options, Copyright (C) 2003-2004,
25 * Paul Clements, SteelEye Technology, Inc.
32 static int scan_assemble(struct supertype
*ss
,
34 struct mddev_ident
*ident
);
35 static int misc_scan(char devmode
, struct context
*c
);
36 static int stop_scan(int verbose
);
37 static int misc_list(struct mddev_dev
*devlist
,
38 struct mddev_ident
*ident
,
40 struct supertype
*ss
, struct context
*c
);
41 const char Name
[] = "mdadm";
43 int main(int argc
, char *argv
[])
51 unsigned long long array_size
= 0;
52 unsigned long long data_offset
= INVALID_SECTORS
;
53 struct mddev_ident ident
;
54 char *configfile
= NULL
;
57 struct mddev_dev
*devlist
= NULL
;
58 struct mddev_dev
**devlistend
= & devlist
;
61 char *symlinks
= NULL
;
62 int grow_continue
= 0;
63 /* autof indicates whether and how to create device node.
64 * bottom 3 bits are style. Rest (when shifted) are number of parts
66 * 1 - don't create (no)
67 * 2 - if is_standard, then create (yes)
68 * 3 - create as 'md' - reject is_standard mdp (md)
69 * 4 - create as 'mdp' - reject is_standard md (mdp)
70 * 5 - default to md if not is_standard (md in config file)
71 * 6 - default to mdp if not is_standard (part, or mdp in config file)
74 .require_homehost
= 1,
80 .bitmap_chunk
= UnSet
,
83 char sys_hostname
[256];
84 char *mailaddr
= NULL
;
90 int spare_sharing
= 1;
91 struct supertype
*ss
= NULL
;
93 char *shortopt
= short_options
;
96 char *remove_path
= NULL
;
97 char *udev_filename
= NULL
;
98 char *dump_directory
= NULL
;
105 srandom(time(0) ^ getpid());
109 ident
.raid_disks
= UnSet
;
110 ident
.super_minor
= UnSet
;
112 ident
.spare_group
= NULL
;
115 ident
.bitmap_fd
= -1;
116 ident
.bitmap_file
= NULL
;
118 ident
.container
= NULL
;
121 while ((option_index
= -1) ,
122 (opt
=getopt_long(argc
, argv
,
123 shortopt
, long_options
,
124 &option_index
)) != -1) {
126 /* firstly, some mode-independent options */
136 fputs(Version
, stderr
);
139 case 'v': c
.verbose
++;
142 case 'q': c
.verbose
--;
146 if (mode
== ASSEMBLE
|| mode
== BUILD
||
147 mode
== CREATE
|| mode
== GROW
||
148 mode
== INCREMENTAL
|| mode
== MANAGE
)
149 break; /* b means bitmap */
154 case 'Y': c
.export
++;
158 if (strcasecmp(optarg
, "<ignore>") == 0)
159 c
.require_homehost
= 0;
165 /* Silently ignore old option */
171 if (asprintf(&c
.prefer
, "/%s/", optarg
) <= 0)
177 fputs(Usage
, stderr
);
180 /* second, figure out the mode.
181 * Some options force the mode. Others
182 * set the mode if it isn't already
188 shortopt
= short_bitmap_options
;
200 case ReAdd
: /* re-add */
204 shortopt
= short_bitmap_options
;
208 case 'A': newmode
= ASSEMBLE
;
209 shortopt
= short_bitmap_auto_options
;
211 case 'B': newmode
= BUILD
;
212 shortopt
= short_bitmap_auto_options
;
214 case 'C': newmode
= CREATE
;
215 shortopt
= short_bitmap_auto_options
;
217 case 'F': newmode
= MONITOR
;
219 case 'G': newmode
= GROW
;
220 shortopt
= short_bitmap_options
;
222 case 'I': newmode
= INCREMENTAL
;
223 shortopt
= short_bitmap_auto_options
;
226 newmode
= AUTODETECT
;
261 if (mode
&& newmode
== mode
) {
262 /* everybody happy ! */
263 } else if (mode
&& newmode
!= mode
) {
266 if (option_index
>= 0)
267 fprintf(stderr
, "--%s", long_options
[option_index
].name
);
269 fprintf(stderr
, "-%c", opt
);
270 fprintf(stderr
, " would set mdadm mode to \"%s\", but it is already set to \"%s\".\n",
271 map_num(modes
, newmode
),
272 map_num(modes
, mode
));
274 } else if (!mode
&& newmode
) {
276 if (mode
== MISC
&& devs_found
) {
277 pr_err("No action given for %s in --misc mode\n",
279 cont_err("Action options must come before device names\n");
283 /* special case of -c --help */
284 if ((opt
== 'c' || opt
== ConfigFile
) &&
285 (strncmp(optarg
, "--h", 3) == 0 ||
286 strncmp(optarg
, "-h", 2) == 0)) {
287 fputs(Help_config
, stdout
);
291 /* If first option is a device, don't force the mode yet */
293 if (devs_found
== 0) {
294 dv
= xmalloc(sizeof(*dv
));
295 dv
->devname
= optarg
;
296 dv
->disposition
= devmode
;
297 dv
->writemostly
= writemostly
;
301 devlistend
= &dv
->next
;
306 /* No mode yet, and this is the second device ... */
307 pr_err("An option must be given to set the mode before a second device\n"
308 " (%s) is listed\n", optarg
);
311 if (option_index
>= 0)
312 pr_err("--%s", long_options
[option_index
].name
);
315 fprintf(stderr
, " does not set the mode, and so cannot be the first option.\n");
319 /* if we just set the mode, then done */
333 /* an undecorated option - must be a device name.
336 if (devs_found
> 0 && devmode
== DetailPlatform
) {
337 pr_err("controller may only be specified once. %s ignored\n",
342 if (devs_found
> 0 && mode
== MANAGE
&& !devmode
) {
343 pr_err("Must give one of -a/-r/-f for subsequent devices at %s\n", optarg
);
346 if (devs_found
> 0 && mode
== GROW
&& !devmode
) {
347 pr_err("Must give -a/--add for devices to add: %s\n", optarg
);
350 dv
= xmalloc(sizeof(*dv
));
351 dv
->devname
= optarg
;
352 dv
->disposition
= devmode
;
353 dv
->writemostly
= writemostly
;
357 devlistend
= &dv
->next
;
363 /* We've got a mode, and opt is now something else which
364 * could depend on the mode */
365 #define O(a,b) ((a<<16)|b)
366 switch (O(mode
,opt
)) {
368 case O(GROW
,ChunkSize
):
370 case O(CREATE
,ChunkSize
):
371 case O(BUILD
,'c'): /* chunk or rounding */
372 case O(BUILD
,ChunkSize
): /* chunk or rounding */
374 pr_err("chunk/rounding may only be specified once. Second value is %s.\n", optarg
);
377 s
.chunk
= parse_size(optarg
);
378 if (s
.chunk
== INVALID_SECTORS
||
379 s
.chunk
< 8 || (s
.chunk
&1)) {
380 pr_err("invalid chunk/rounding value: %s\n",
384 /* Convert sectors to K */
388 case O(INCREMENTAL
, 'e'):
390 case O(ASSEMBLE
,'e'):
391 case O(MISC
,'e'): /* set metadata (superblock) information */
393 pr_err("metadata information already given\n");
396 for(i
=0; !ss
&& superlist
[i
]; i
++)
397 ss
= superlist
[i
]->match_metadata_desc(optarg
);
400 pr_err("unrecognised metadata identifier: %s\n", optarg
);
406 case O(MANAGE
,WriteMostly
):
408 case O(BUILD
,WriteMostly
):
410 case O(CREATE
,WriteMostly
):
411 /* set write-mostly for following devices */
416 /* clear write-mostly for following devices */
422 case O(BUILD
,'z'): /* size */
424 pr_err("size may only be specified once. Second value is %s.\n", optarg
);
427 if (strcmp(optarg
, "max") == 0)
430 s
.size
= parse_size(optarg
);
431 if (s
.size
== INVALID_SECTORS
|| s
.size
< 8) {
432 pr_err("invalid size: %s\n", optarg
);
435 /* convert sectors to K */
440 case O(GROW
,'Z'): /* array size */
441 if (array_size
> 0) {
442 pr_err("array-size may only be specified once. Second value is %s.\n", optarg
);
445 if (strcmp(optarg
, "max") == 0)
446 array_size
= MAX_SIZE
;
448 array_size
= parse_size(optarg
);
449 if (array_size
== 0 ||
450 array_size
== INVALID_SECTORS
) {
451 pr_err("invalid array size: %s\n",
458 case O(CREATE
,DataOffset
):
459 case O(GROW
,DataOffset
):
460 if (data_offset
!= INVALID_SECTORS
) {
461 pr_err("data-offset may only be specified one. Second value is %s.\n", optarg
);
464 if (mode
== CREATE
&& strcmp(optarg
, "variable") == 0)
465 data_offset
= VARIABLE_OFFSET
;
467 data_offset
= parse_size(optarg
);
468 if (data_offset
== INVALID_SECTORS
) {
469 pr_err("invalid data-offset: %s\n",
477 case O(BUILD
,'l'): /* set raid level*/
478 if (s
.level
!= UnSet
) {
479 pr_err("raid level may only be set once. Second value is %s.\n", optarg
);
482 s
.level
= map_name(pers
, optarg
);
483 if (s
.level
== UnSet
) {
484 pr_err("invalid raid level: %s\n",
488 if (s
.level
!= 0 && s
.level
!= LEVEL_LINEAR
&&
489 s
.level
!= 1 && s
.level
!= LEVEL_MULTIPATH
&&
490 s
.level
!= LEVEL_FAULTY
&& s
.level
!= 10 &&
492 pr_err("Raid level %s not permitted with --build.\n",
496 if (s
.sparedisks
> 0 && s
.level
< 1 && s
.level
>= -1) {
497 pr_err("raid level %s is incompatible with spare-devices setting.\n",
501 ident
.level
= s
.level
;
504 case O(GROW
, 'p'): /* new layout */
505 case O(GROW
, Layout
):
507 pr_err("layout may only be sent once. Second value was %s\n", optarg
);
510 s
.layout_str
= optarg
;
511 /* 'Grow' will parse the value */
514 case O(CREATE
,'p'): /* raid5 layout */
515 case O(CREATE
,Layout
):
516 case O(BUILD
,'p'): /* faulty layout */
517 case O(BUILD
,Layout
):
518 if (s
.layout
!= UnSet
) {
519 pr_err("layout may only be sent once. Second value was %s\n", optarg
);
524 pr_err("layout not meaningful for %s arrays.\n",
525 map_num(pers
, s
.level
));
528 pr_err("raid level must be given before layout.\n");
532 s
.layout
= map_name(r5layout
, optarg
);
533 if (s
.layout
==UnSet
) {
534 pr_err("layout %s not understood for raid5.\n",
540 s
.layout
= map_name(r6layout
, optarg
);
541 if (s
.layout
==UnSet
) {
542 pr_err("layout %s not understood for raid6.\n",
549 s
.layout
= parse_layout_10(optarg
);
551 pr_err("layout for raid10 must be 'nNN', 'oNN' or 'fNN' where NN is a number, not %s\n", optarg
);
559 s
.layout
= parse_layout_faulty(optarg
);
560 if (s
.layout
== -1) {
561 pr_err("layout %s not understood for faulty.\n",
569 case O(CREATE
,AssumeClean
):
570 case O(BUILD
,AssumeClean
): /* assume clean */
571 case O(GROW
,AssumeClean
):
577 case O(BUILD
,'n'): /* number of raid disks */
579 pr_err("raid-devices set twice: %d and %s\n",
580 s
.raiddisks
, optarg
);
583 s
.raiddisks
= parse_num(optarg
);
584 if (s
.raiddisks
<= 0) {
585 pr_err("invalid number of raid devices: %s\n",
589 ident
.raid_disks
= s
.raiddisks
;
591 case O(ASSEMBLE
, Nodes
):
593 case O(CREATE
, Nodes
):
594 c
.nodes
= parse_num(optarg
);
596 pr_err("invalid number for the number of cluster nodes: %s\n",
601 case O(CREATE
, ClusterName
):
602 case O(ASSEMBLE
, ClusterName
):
603 c
.homecluster
= optarg
;
604 if (strlen(c
.homecluster
) > 64) {
605 pr_err("Cluster name too big.\n");
609 case O(CREATE
,'x'): /* number of spare (eXtra) disks */
611 pr_err("spare-devices set twice: %d and %s\n",
612 s
.sparedisks
, optarg
);
615 if (s
.level
!= UnSet
&& s
.level
<= 0 && s
.level
>= -1) {
616 pr_err("spare-devices setting is incompatible with raid level %d\n",
620 s
.sparedisks
= parse_num(optarg
);
621 if (s
.sparedisks
< 0) {
622 pr_err("invalid number of spare-devices: %s\n",
632 case O(INCREMENTAL
,'a'):
633 case O(INCREMENTAL
,Auto
):
634 case O(ASSEMBLE
,'a'):
635 case O(ASSEMBLE
,Auto
): /* auto-creation of device node */
636 c
.autof
= parse_auto(optarg
, "--auto flag", 0);
639 case O(CREATE
,Symlinks
):
640 case O(BUILD
,Symlinks
):
641 case O(ASSEMBLE
,Symlinks
): /* auto creation of symlinks in /dev to /dev/md */
645 case O(BUILD
,'f'): /* force honouring '-n 1' */
646 case O(BUILD
,Force
): /* force honouring '-n 1' */
647 case O(GROW
,'f'): /* ditto */
648 case O(GROW
,Force
): /* ditto */
649 case O(CREATE
,'f'): /* force honouring of device list */
650 case O(CREATE
,Force
): /* force honouring of device list */
651 case O(ASSEMBLE
,'f'): /* force assembly */
652 case O(ASSEMBLE
,Force
): /* force assembly */
653 case O(MISC
,'f'): /* force zero */
654 case O(MISC
,Force
): /* force zero */
655 case O(MANAGE
,Force
): /* add device which is too large */
658 /* now for the Assemble options */
659 case O(ASSEMBLE
, FreezeReshape
): /* Freeze reshape during
661 case O(INCREMENTAL
, FreezeReshape
):
662 c
.freeze_reshape
= 1;
664 case O(CREATE
,'u'): /* uuid of array */
665 case O(ASSEMBLE
,'u'): /* uuid of array */
666 if (ident
.uuid_set
) {
667 pr_err("uuid cannot be set twice. Second value %s.\n", optarg
);
670 if (parse_uuid(optarg
, ident
.uuid
))
673 pr_err("Bad uuid: %s\n", optarg
);
679 case O(ASSEMBLE
,'N'):
682 pr_err("name cannot be set twice. Second value %s.\n", optarg
);
685 if (mode
== MISC
&& !c
.subarray
) {
686 pr_err("-N/--name only valid with --update-subarray in misc mode\n");
689 if (strlen(optarg
) > 32) {
690 pr_err("name '%s' is too long, 32 chars max.\n",
694 strcpy(ident
.name
, optarg
);
697 case O(ASSEMBLE
,'m'): /* super-minor for array */
698 case O(ASSEMBLE
,SuperMinor
):
699 if (ident
.super_minor
!= UnSet
) {
700 pr_err("super-minor cannot be set twice. Second value: %s.\n", optarg
);
703 if (strcmp(optarg
, "dev") == 0)
704 ident
.super_minor
= -2;
706 ident
.super_minor
= parse_num(optarg
);
707 if (ident
.super_minor
< 0) {
708 pr_err("Bad super-minor number: %s.\n", optarg
);
714 case O(ASSEMBLE
,'o'):
720 case O(ASSEMBLE
,'U'): /* update the superblock */
723 pr_err("Can only update one aspect of superblock, both %s and %s given.\n",
727 if (mode
== MISC
&& !c
.subarray
) {
728 pr_err("Only subarrays can be updated in misc mode\n");
732 if (strcmp(c
.update
, "sparc2.2") == 0)
734 if (strcmp(c
.update
, "super-minor") == 0)
736 if (strcmp(c
.update
, "summaries") == 0)
738 if (strcmp(c
.update
, "resync") == 0)
740 if (strcmp(c
.update
, "uuid") == 0)
742 if (strcmp(c
.update
, "name") == 0)
744 if (strcmp(c
.update
, "homehost") == 0)
746 if (strcmp(c
.update
, "home-cluster") == 0)
748 if (strcmp(c
.update
, "nodes") == 0)
750 if (strcmp(c
.update
, "devicesize") == 0)
752 if (strcmp(c
.update
, "no-bitmap") == 0)
754 if (strcmp(c
.update
, "bbl") == 0)
756 if (strcmp(c
.update
, "no-bbl") == 0)
758 if (strcmp(c
.update
, "force-no-bbl") == 0)
760 if (strcmp(c
.update
, "metadata") == 0)
762 if (strcmp(c
.update
, "revert-reshape") == 0)
764 if (strcmp(c
.update
, "byteorder")==0) {
766 pr_err("must not set metadata type with --update=byteorder.\n");
769 for(i
=0; !ss
&& superlist
[i
]; i
++)
770 ss
= superlist
[i
]->match_metadata_desc(
773 pr_err("INTERNAL ERROR cannot find 0.swap\n");
779 if (strcmp(c
.update
,"?") == 0 ||
780 strcmp(c
.update
, "help") == 0) {
782 fprintf(outf
, "%s: ", Name
);
786 "%s: '--update=%s' is invalid. ",
789 fprintf(outf
, "Valid --update options are:\n"
790 " 'sparc2.2', 'super-minor', 'uuid', 'name', 'nodes', 'resync',\n"
791 " 'summaries', 'homehost', 'home-cluster', 'byteorder', 'devicesize',\n"
792 " 'no-bitmap', 'metadata', 'revert-reshape'\n"
793 " 'bbl', 'no-bbl', 'force-no-bbl'\n"
795 exit(outf
== stdout
? 0 : 2);
798 /* update=devicesize is allowed with --re-add */
799 if (devmode
!= 'A') {
800 pr_err("--update in Manage mode only allowed with --re-add.\n");
804 pr_err("Can only update one aspect of superblock, both %s and %s given.\n",
809 if (strcmp(c
.update
, "devicesize") != 0 &&
810 strcmp(c
.update
, "bbl") != 0 &&
811 strcmp(c
.update
, "force-no-bbl") != 0 &&
812 strcmp(c
.update
, "no-bbl") != 0) {
813 pr_err("only 'devicesize', 'bbl', 'no-bbl', and 'force-no-bbl' can be updated with --re-add\n");
818 case O(INCREMENTAL
,NoDegraded
):
819 pr_err("--no-degraded is deprecated in Incremental mode\n");
820 case O(ASSEMBLE
,NoDegraded
): /* --no-degraded */
821 c
.runstop
= -1; /* --stop isn't allowed for --assemble,
822 * so we overload slightly */
825 case O(ASSEMBLE
,'c'):
826 case O(ASSEMBLE
,ConfigFile
):
827 case O(INCREMENTAL
, 'c'):
828 case O(INCREMENTAL
, ConfigFile
):
830 case O(MISC
, ConfigFile
):
832 case O(MONITOR
,ConfigFile
):
833 case O(CREATE
,ConfigFile
):
835 pr_err("configfile cannot be set twice. Second value is %s.\n", optarg
);
839 set_conffile(configfile
);
840 /* FIXME possibly check that config file exists. Even parse it */
842 case O(ASSEMBLE
,'s'): /* scan */
845 case O(INCREMENTAL
,'s'):
849 case O(MONITOR
,'m'): /* mail address */
850 case O(MONITOR
,EMail
):
852 pr_err("only specify one mailaddress. %s ignored.\n",
858 case O(MONITOR
,'p'): /* alert program */
859 case O(MONITOR
,ProgramOpt
): /* alert program */
861 pr_err("only specify one alter program. %s ignored.\n",
867 case O(MONITOR
,'r'): /* rebuild increments */
868 case O(MONITOR
,Increment
):
869 increments
= atoi(optarg
);
870 if (increments
> 99 || increments
< 1) {
871 pr_err("please specify positive integer between 1 and 99 as rebuild increments.\n");
876 case O(MONITOR
,'d'): /* delay in seconds */
878 case O(BUILD
,'d'): /* delay for bitmap updates */
881 pr_err("only specify delay once. %s ignored.\n",
884 c
.delay
= parse_num(optarg
);
886 pr_err("invalid delay: %s\n",
892 case O(MONITOR
,'f'): /* daemonise */
893 case O(MONITOR
,Fork
):
896 case O(MONITOR
,'i'): /* pid */
898 pr_err("only specify one pid file. %s ignored.\n",
903 case O(MONITOR
,'1'): /* oneshot */
907 case O(MONITOR
,'t'): /* test */
910 case O(MONITOR
,'y'): /* log messages to syslog */
911 openlog("mdadm", LOG_PID
, SYSLOG_FACILITY
);
914 case O(MONITOR
, NoSharing
):
918 /* now the general management options. Some are applicable
919 * to other modes. None have arguments.
924 case O(MANAGE
,Add
): /* add a drive */
927 case O(MANAGE
,AddSpare
): /* add drive - never re-add */
930 case O(MANAGE
,AddJournal
): /* add journal */
931 if (s
.journaldisks
&& (s
.level
< 4 || s
.level
> 6)) {
932 pr_err("--add-journal is only supported for RAID level 4/5/6.\n");
937 case O(MANAGE
,ReAdd
):
940 case O(MANAGE
,'r'): /* remove a drive */
941 case O(MANAGE
,Remove
):
944 case O(MANAGE
,'f'): /* set faulty */
946 case O(INCREMENTAL
,'f'):
947 case O(INCREMENTAL
,Remove
):
948 case O(INCREMENTAL
,Fail
): /* r for incremental is taken, use f
949 * even though we will both fail and
950 * remove the device */
953 case O(MANAGE
, ClusterConfirm
):
956 case O(MANAGE
,Replace
):
957 /* Mark these devices for replacement */
961 /* These are the replacements to use */
962 if (devmode
!= 'R') {
963 pr_err("--with must follow --replace\n");
968 case O(INCREMENTAL
,'R'):
970 case O(ASSEMBLE
,'R'):
972 case O(CREATE
,'R'): /* Run the array */
974 pr_err("Cannot both Stop and Run an array\n");
981 pr_err("Cannot both Run and Stop an array\n");
993 case O(MISC
,KillOpt
):
997 case O(MISC
, ExamineBB
):
1001 case O(MISC
, WaitOpt
):
1002 case O(MISC
, Waitclean
):
1003 case O(MISC
, DetailPlatform
):
1004 case O(MISC
, KillSubarray
):
1005 case O(MISC
, UpdateSubarray
):
1007 case O(MISC
, Restore
):
1008 case O(MISC
,Action
):
1009 if (opt
== KillSubarray
|| opt
== UpdateSubarray
) {
1011 pr_err("subarray can only be specified once\n");
1014 c
.subarray
= optarg
;
1016 if (opt
== Action
) {
1018 pr_err("Only one --action can be specified\n");
1021 if (strcmp(optarg
, "idle") == 0 ||
1022 strcmp(optarg
, "frozen") == 0 ||
1023 strcmp(optarg
, "check") == 0 ||
1024 strcmp(optarg
, "repair") == 0)
1027 pr_err("action must be one of idle, frozen, check, repair\n");
1031 if (devmode
&& devmode
!= opt
&&
1033 (opt
== 'E' && devmode
!= 'Q'))) {
1034 pr_err("--examine/-E cannot be given with ");
1035 if (devmode
== 'E') {
1036 if (option_index
>= 0)
1037 fprintf(stderr
, "--%s\n",
1038 long_options
[option_index
].name
);
1040 fprintf(stderr
, "-%c\n", opt
);
1041 } else if (isalpha(devmode
))
1042 fprintf(stderr
, "-%c\n", devmode
);
1044 fprintf(stderr
, "previous option\n");
1048 if (opt
== Dump
|| opt
== Restore
) {
1049 if (dump_directory
!= NULL
) {
1050 pr_err("dump/restore directory specified twice: %s and %s\n",
1051 dump_directory
, optarg
);
1054 dump_directory
= optarg
;
1057 case O(MISC
, UdevRules
):
1058 if (devmode
&& devmode
!= opt
) {
1059 pr_err("--udev-rules must be the only option.\n");
1062 pr_err("only specify one udev rule filename. %s ignored.\n",
1065 udev_filename
= optarg
;
1073 case O(MISC
, Sparc22
):
1074 if (devmode
!= 'E') {
1075 pr_err("--sparc2.2 only allowed with --examine\n");
1081 case O(ASSEMBLE
,'b'): /* here we simply set the bitmap file */
1082 case O(ASSEMBLE
,Bitmap
):
1084 pr_err("bitmap file needed with -b in --assemble mode\n");
1087 if (strcmp(optarg
, "internal") == 0) {
1088 pr_err("there is no need to specify --bitmap when assembling arrays with internal bitmaps\n");
1091 bitmap_fd
= open(optarg
, O_RDWR
);
1092 if (!*optarg
|| bitmap_fd
< 0) {
1093 pr_err("cannot open bitmap file %s: %s\n", optarg
, strerror(errno
));
1096 ident
.bitmap_fd
= bitmap_fd
; /* for Assemble */
1099 case O(ASSEMBLE
, BackupFile
):
1100 case O(GROW
, BackupFile
):
1101 /* Specify a file into which grow might place a backup,
1102 * or from which assemble might recover a backup
1104 if (c
.backup_file
) {
1105 pr_err("backup file already specified, rejecting %s\n", optarg
);
1108 c
.backup_file
= optarg
;
1111 case O(GROW
, Continue
):
1112 /* Continue interrupted grow
1116 case O(ASSEMBLE
, InvalidBackup
):
1117 /* Acknowledge that the backupfile is invalid, but ask
1118 * to continue anyway
1120 c
.invalid_backup
= 1;
1124 case O(BUILD
,Bitmap
):
1126 case O(CREATE
,Bitmap
): /* here we create the bitmap */
1128 case O(GROW
,Bitmap
):
1129 if (strcmp(optarg
, "internal") == 0 ||
1130 strcmp(optarg
, "none") == 0 ||
1131 strchr(optarg
, '/') != NULL
) {
1132 s
.bitmap_file
= optarg
;
1135 if (strcmp(optarg
, "clustered") == 0) {
1136 s
.bitmap_file
= optarg
;
1137 /* Set the default number of cluster nodes
1138 * to 4 if not already set by user
1145 pr_err("bitmap file must contain a '/', or be 'internal', or 'none'\n"
1146 " not '%s'\n", optarg
);
1149 case O(GROW
,BitmapChunk
):
1150 case O(BUILD
,BitmapChunk
):
1151 case O(CREATE
,BitmapChunk
): /* bitmap chunksize */
1152 s
.bitmap_chunk
= parse_size(optarg
);
1153 if (s
.bitmap_chunk
== 0 ||
1154 s
.bitmap_chunk
== INVALID_SECTORS
||
1155 s
.bitmap_chunk
& (s
.bitmap_chunk
- 1)) {
1156 pr_err("invalid bitmap chunksize: %s\n",
1160 s
.bitmap_chunk
= s
.bitmap_chunk
* 512;
1163 case O(GROW
, WriteBehind
):
1164 case O(BUILD
, WriteBehind
):
1165 case O(CREATE
, WriteBehind
): /* write-behind mode */
1166 s
.write_behind
= DEFAULT_MAX_WRITE_BEHIND
;
1168 s
.write_behind
= parse_num(optarg
);
1169 if (s
.write_behind
< 0 ||
1170 s
.write_behind
> 16383) {
1171 pr_err("Invalid value for maximum outstanding write-behind writes: %s.\n\tMust be between 0 and 16383.\n", optarg
);
1177 case O(INCREMENTAL
, 'r'):
1178 case O(INCREMENTAL
, RebuildMapOpt
):
1181 case O(INCREMENTAL
, IncrementalPath
):
1182 remove_path
= optarg
;
1184 case O(CREATE
, WriteJournal
):
1185 if (s
.journaldisks
) {
1186 pr_err("Please specify only one journal device for the array.\n");
1187 pr_err("Ignoring --write-journal %s...\n", optarg
);
1190 dv
= xmalloc(sizeof(*dv
));
1191 dv
->devname
= optarg
;
1192 dv
->disposition
= 'j'; /* WriteJournal */
1196 devlistend
= &dv
->next
;
1202 /* We have now processed all the valid options. Anything else is
1205 if (option_index
> 0)
1206 pr_err(":option --%s not valid in %s mode\n",
1207 long_options
[option_index
].name
,
1208 map_num(modes
, mode
));
1210 pr_err("option -%c not valid in %s mode\n",
1211 opt
, map_num(modes
, mode
));
1218 if (print_help
== 2)
1219 help_text
= OptionHelp
;
1221 help_text
= mode_help
[mode
];
1222 if (help_text
== NULL
)
1224 fputs(help_text
,stdout
);
1228 if (s
.journaldisks
&& (s
.level
< 4 || s
.level
> 6)) {
1229 pr_err("--write-journal is only supported for RAID level 4/5/6.\n");
1233 if (!mode
&& devs_found
) {
1236 if (devlist
->disposition
== 0)
1237 devlist
->disposition
= devmode
;
1240 fputs(Usage
, stderr
);
1245 struct createinfo
*ci
= conf_get_create_info();
1247 if (strcasecmp(symlinks
, "yes") == 0)
1249 else if (strcasecmp(symlinks
, "no") == 0)
1252 pr_err("option --symlinks must be 'no' or 'yes'\n");
1256 /* Ok, got the option parsing out of the way
1257 * hopefully it's mostly right but there might be some stuff
1260 * That is mosty checked in the per-mode stuff but...
1262 * For @,B,C and A without -s, the first device listed must be
1263 * an md device. We check that here and open it.
1266 if (mode
== MANAGE
|| mode
== BUILD
|| mode
== CREATE
||
1267 mode
== GROW
|| (mode
== ASSEMBLE
&& ! c
.scan
)) {
1268 if (devs_found
< 1) {
1269 pr_err("an md device must be given in this mode\n");
1272 if ((int)ident
.super_minor
== -2 && c
.autof
) {
1273 pr_err("--super-minor=dev is incompatible with --auto\n");
1276 if (mode
== MANAGE
|| mode
== GROW
) {
1277 mdfd
= open_mddev(devlist
->devname
, 1);
1281 /* non-existent device is OK */
1282 mdfd
= open_mddev(devlist
->devname
, 0);
1284 pr_err("device %s exists but is not an md array.\n", devlist
->devname
);
1287 if ((int)ident
.super_minor
== -2) {
1290 pr_err("--super-minor=dev given, and listed device %s doesn't exist.\n",
1295 ident
.super_minor
= minor(stb
.st_rdev
);
1297 if (mdfd
>= 0 && mode
!= MANAGE
&& mode
!= GROW
) {
1298 /* We don't really want this open yet, we just might
1299 * have wanted to check some things
1307 if (s
.raiddisks
== 1 && !c
.force
&& s
.level
!= LEVEL_FAULTY
) {
1308 pr_err("'1' is an unusual number of drives for an array, so it is probably\n"
1309 " a mistake. If you really mean it you will need to specify --force before\n"
1310 " setting the number of drives.\n");
1315 if (c
.homehost
== NULL
&& c
.require_homehost
)
1316 c
.homehost
= conf_get_homehost(&c
.require_homehost
);
1317 if (c
.homehost
== NULL
|| strcasecmp(c
.homehost
, "<system>") == 0) {
1318 if (gethostname(sys_hostname
, sizeof(sys_hostname
)) == 0) {
1319 sys_hostname
[sizeof(sys_hostname
)-1] = 0;
1320 c
.homehost
= sys_hostname
;
1324 (!c
.homehost
[0] || strcasecmp(c
.homehost
, "<none>") == 0)) {
1326 c
.require_homehost
= 0;
1331 set_hooks(); /* set hooks from libs */
1333 if (c
.homecluster
== NULL
&& (c
.nodes
> 0)) {
1334 c
.homecluster
= conf_get_homecluster();
1335 if (c
.homecluster
== NULL
)
1336 rv
= get_cluster_name(&c
.homecluster
);
1338 pr_err("The md can't get cluster name\n");
1343 if (c
.backup_file
&& data_offset
!= INVALID_SECTORS
) {
1344 pr_err("--backup-file and --data-offset are incompatible\n");
1348 if ((mode
== MISC
&& devmode
== 'E') ||
1349 (mode
== MONITOR
&& spare_sharing
== 0))
1350 /* Anyone may try this */;
1351 else if (geteuid() != 0) {
1352 pr_err("must be super-user to perform this action\n");
1356 ident
.autof
= c
.autof
;
1358 if (c
.scan
&& c
.verbose
< 2)
1359 /* --scan implied --brief unless -vv */
1364 /* readonly, add/remove, readwrite, runstop */
1366 rv
= Manage_ro(devlist
->devname
, mdfd
, c
.readonly
);
1367 if (!rv
&& devs_found
>1)
1368 rv
= Manage_subdevs(devlist
->devname
, mdfd
,
1369 devlist
->next
, c
.verbose
, c
.test
,
1371 if (!rv
&& c
.readonly
< 0)
1372 rv
= Manage_ro(devlist
->devname
, mdfd
, c
.readonly
);
1373 if (!rv
&& c
.runstop
> 0)
1374 rv
= Manage_run(devlist
->devname
, mdfd
, &c
);
1375 if (!rv
&& c
.runstop
< 0)
1376 rv
= Manage_stop(devlist
->devname
, mdfd
, c
.verbose
, 0);
1379 if (devs_found
== 1 && ident
.uuid_set
== 0 &&
1380 ident
.super_minor
== UnSet
&& ident
.name
[0] == 0 &&
1382 /* Only a device has been given, so get details from config file */
1383 struct mddev_ident
*array_ident
= conf_get_ident(devlist
->devname
);
1384 if (array_ident
== NULL
) {
1385 pr_err("%s not identified in config file.\n",
1391 if (array_ident
->autof
== 0)
1392 array_ident
->autof
= c
.autof
;
1393 rv
|= Assemble(ss
, devlist
->devname
, array_ident
,
1397 rv
= Assemble(ss
, devlist
->devname
, &ident
,
1399 else if (devs_found
> 0) {
1400 if (c
.update
&& devs_found
> 1) {
1401 pr_err("can only update a single array at a time\n");
1404 if (c
.backup_file
&& devs_found
> 1) {
1405 pr_err("can only assemble a single array when providing a backup file.\n");
1408 for (dv
= devlist
; dv
; dv
=dv
->next
) {
1409 struct mddev_ident
*array_ident
= conf_get_ident(dv
->devname
);
1410 if (array_ident
== NULL
) {
1411 pr_err("%s not identified in config file.\n",
1416 if (array_ident
->autof
== 0)
1417 array_ident
->autof
= c
.autof
;
1418 rv
|= Assemble(ss
, dv
->devname
, array_ident
,
1423 pr_err("--update not meaningful with a --scan assembly.\n");
1426 if (c
.backup_file
) {
1427 pr_err("--backup_file not meaningful with a --scan assembly.\n");
1430 rv
= scan_assemble(ss
, &c
, &ident
);
1436 c
.delay
= DEFAULT_BITMAP_DELAY
;
1437 if (s
.write_behind
&& !s
.bitmap_file
) {
1438 pr_err("write-behind mode requires a bitmap.\n");
1442 if (s
.raiddisks
== 0) {
1443 pr_err("no raid-devices specified.\n");
1448 if (s
.bitmap_file
) {
1449 if (strcmp(s
.bitmap_file
, "internal") == 0 ||
1450 strcmp(s
.bitmap_file
, "clustered") == 0) {
1451 pr_err("'internal' and 'clustered' bitmaps not supported with --build\n");
1456 rv
= Build(devlist
->devname
, devlist
->next
, &s
, &c
);
1460 c
.delay
= DEFAULT_BITMAP_DELAY
;
1463 if (!s
.bitmap_file
||
1464 strcmp(s
.bitmap_file
, "clustered") != 0) {
1465 pr_err("--nodes argument only compatible with --bitmap=clustered\n");
1471 pr_err("--bitmap=clustered is currently supported with RAID mirror only\n");
1477 if (s
.write_behind
&& !s
.bitmap_file
) {
1478 pr_err("write-behind mode requires a bitmap.\n");
1482 if (s
.raiddisks
== 0) {
1483 pr_err("no raid-devices specified.\n");
1488 rv
= Create(ss
, devlist
->devname
,
1489 ident
.name
, ident
.uuid_set
? ident
.uuid
: NULL
,
1490 devs_found
-1, devlist
->next
,
1491 &s
, &c
, data_offset
);
1494 if (devmode
== 'E') {
1495 if (devlist
== NULL
&& !c
.scan
) {
1496 pr_err("No devices to examine\n");
1499 if (devlist
== NULL
)
1500 devlist
= conf_get_devs();
1501 if (devlist
== NULL
) {
1502 pr_err("No devices listed in %s\n", configfile
?configfile
:DefaultConfFile
);
1505 rv
= Examine(devlist
, &c
, ss
);
1506 } else if (devmode
== DetailPlatform
) {
1507 rv
= Detail_Platform(ss
? ss
->ss
: NULL
, ss
? c
.scan
: 1,
1508 c
.verbose
, c
.export
,
1509 devlist
? devlist
->devname
: NULL
);
1510 } else if (devlist
== NULL
) {
1511 if (devmode
== 'S' && c
.scan
)
1512 rv
= stop_scan(c
.verbose
);
1513 else if ((devmode
== 'D' || devmode
== Waitclean
) &&
1515 rv
= misc_scan(devmode
, &c
);
1516 else if (devmode
== UdevRules
)
1517 rv
= Write_rules(udev_filename
);
1519 pr_err("No devices given.\n");
1523 rv
= misc_list(devlist
, &ident
, dump_directory
, ss
, &c
);
1526 if (!devlist
&& !c
.scan
) {
1527 pr_err("Cannot monitor: need --scan or at least one device\n");
1531 if (pidfile
&& !daemonise
) {
1532 pr_err("Cannot write a pid file when not in daemon mode\n");
1537 if (get_linux_version() > 2006016)
1538 /* mdstat responds to poll */
1543 rv
= Monitor(devlist
, mailaddr
, program
,
1544 &c
, daemonise
, oneshot
,
1545 dosyslog
, pidfile
, increments
,
1550 if (array_size
> 0) {
1551 /* alway impose array size first, independent of
1553 * Do not allow level or raid_disks changes at the
1554 * same time as that can be irreversibly destructive.
1558 if (s
.raiddisks
|| s
.level
!= UnSet
) {
1559 pr_err("cannot change array size in same operation as changing raiddisks or level.\n"
1560 " Change size first, then check that data is still intact.\n");
1564 sysfs_init(&sra
, mdfd
, NULL
);
1565 if (array_size
== MAX_SIZE
)
1566 err
= sysfs_set_str(&sra
, NULL
, "array_size", "default");
1568 err
= sysfs_set_num(&sra
, NULL
, "array_size", array_size
/ 2);
1571 pr_err("--array-size setting is too large.\n");
1573 pr_err("current kernel does not support setting --array-size\n");
1578 if (devs_found
> 1 && s
.raiddisks
== 0 && s
.level
== UnSet
) {
1580 if (s
.size
> 0 || s
.chunk
||
1581 s
.layout_str
|| s
.bitmap_file
) {
1582 pr_err("--add cannot be used with other geometry changes in --grow mode\n");
1586 for (dv
=devlist
->next
; dv
; dv
=dv
->next
) {
1587 rv
= Grow_Add_device(devlist
->devname
, mdfd
,
1592 } else if (s
.bitmap_file
) {
1593 if (s
.size
> 0 || s
.raiddisks
|| s
.chunk
||
1594 s
.layout_str
|| devs_found
> 1) {
1595 pr_err("--bitmap changes cannot be used with other geometry changes in --grow mode\n");
1600 c
.delay
= DEFAULT_BITMAP_DELAY
;
1601 rv
= Grow_addbitmap(devlist
->devname
, mdfd
, &c
, &s
);
1602 } else if (grow_continue
)
1603 rv
= Grow_continue_command(devlist
->devname
,
1604 mdfd
, c
.backup_file
,
1606 else if (s
.size
> 0 || s
.raiddisks
|| s
.layout_str
||
1607 s
.chunk
!= 0 || s
.level
!= UnSet
||
1608 data_offset
!= INVALID_SECTORS
) {
1609 rv
= Grow_reshape(devlist
->devname
, mdfd
,
1611 data_offset
, &c
, &s
);
1612 } else if (array_size
== 0)
1613 pr_err("no changes to --grow\n");
1622 pr_err("In --incremental mode, a device cannot be given with --scan.\n");
1625 if (c
.runstop
<= 0) {
1626 pr_err("--incremental --scan meaningless without --run.\n");
1629 if (devmode
== 'f') {
1630 pr_err("--incremental --scan --fail not supported.\n");
1633 rv
= IncrementalScan(&c
, NULL
);
1636 if (!rebuild_map
&& !c
.scan
) {
1637 pr_err("--incremental requires a device.\n");
1642 if (devmode
== 'f') {
1643 if (devlist
->next
) {
1644 pr_err("'--incremental --fail' can only handle one device.\n");
1648 rv
= IncrementalRemove(devlist
->devname
, remove_path
,
1651 rv
= Incremental(devlist
, &c
, ss
);
1660 static int scan_assemble(struct supertype
*ss
,
1662 struct mddev_ident
*ident
)
1664 struct mddev_ident
*a
, *array_list
= conf_get_ident(NULL
);
1665 struct mddev_dev
*devlist
= conf_get_devs();
1666 struct map_ent
*map
= NULL
;
1669 int failures
, successes
;
1671 if (conf_verify_devnames(array_list
)) {
1672 pr_err("Duplicate MD device names in conf file were found.\n");
1675 if (devlist
== NULL
) {
1676 pr_err("No devices listed in conf file were found.\n");
1679 for (a
= array_list
; a
; a
= a
->next
) {
1682 a
->autof
= c
->autof
;
1685 pr_err("failed to get exclusive lock on mapfile\n");
1690 for (a
= array_list
; a
; a
= a
->next
) {
1695 strcasecmp(a
->devname
, "<ignore>") == 0)
1698 r
= Assemble(ss
, a
->devname
,
1708 } while (failures
&& successes
);
1709 if (c
->homehost
&& cnt
== 0) {
1710 /* Maybe we can auto-assemble something.
1711 * Repeatedly call Assemble in auto-assemble mode
1716 ident
->autof
= c
->autof
;
1718 struct mddev_dev
*devlist
= conf_get_devs();
1721 rv2
= Assemble(ss
, NULL
,
1729 /* Incase there are stacked devices, we need to go around again */
1731 if (cnt
== 0 && rv
== 0) {
1732 pr_err("No arrays found in config file or automatically\n");
1736 } else if (cnt
== 0 && rv
== 0) {
1737 pr_err("No arrays found in config file\n");
1744 static int misc_scan(char devmode
, struct context
*c
)
1746 /* apply --detail or --wait-clean to
1747 * all devices in /proc/mdstat
1749 struct mdstat_ent
*ms
= mdstat_read(0, 1);
1750 struct mdstat_ent
*e
;
1751 struct map_ent
*map
= NULL
;
1755 for (members
= 0; members
<= 1; members
++) {
1756 for (e
=ms
; e
; e
=e
->next
) {
1760 int member
= e
->metadata_version
&&
1761 strncmp(e
->metadata_version
,
1762 "external:/", 10) == 0;
1763 if (members
!= member
)
1765 me
= map_by_devnm(&map
, e
->devnm
);
1767 && strcmp(me
->path
, "/unknown") != 0)
1769 if (name
== NULL
|| stat(name
, &stb
) != 0)
1770 name
= get_md_name(e
->devnm
);
1773 pr_err("cannot find device file for %s\n",
1778 rv
|= Detail(name
, c
);
1780 rv
|= WaitClean(name
, -1, c
->verbose
);
1788 static int stop_scan(int verbose
)
1790 /* apply --stop to all devices in /proc/mdstat */
1791 /* Due to possible stacking of devices, repeat until
1792 * nothing more can be stopped
1794 int progress
=1, err
;
1798 struct mdstat_ent
*ms
= mdstat_read(0, 0);
1799 struct mdstat_ent
*e
;
1801 if (!progress
) last
= 1;
1802 progress
= 0; err
= 0;
1803 for (e
=ms
; e
; e
=e
->next
) {
1804 char *name
= get_md_name(e
->devnm
);
1808 pr_err("cannot find device file for %s\n",
1812 mdfd
= open_mddev(name
, 1);
1814 if (Manage_stop(name
, mdfd
, verbose
, !last
))
1824 } while (!last
&& err
);
1830 static int misc_list(struct mddev_dev
*devlist
,
1831 struct mddev_ident
*ident
,
1832 char *dump_directory
,
1833 struct supertype
*ss
, struct context
*c
)
1835 struct mddev_dev
*dv
;
1838 for (dv
=devlist
; dv
; dv
=(rv
& 16) ? NULL
: dv
->next
) {
1841 switch(dv
->disposition
) {
1843 rv
|= Detail(dv
->devname
, c
);
1845 case KillOpt
: /* Zero superblock */
1847 rv
|= Kill(dv
->devname
, ss
, c
->force
, c
->verbose
,0);
1851 rv
|= Kill(dv
->devname
, NULL
, c
->force
, v
, 0);
1858 rv
|= Query(dv
->devname
); continue;
1860 rv
|= ExamineBitmap(dv
->devname
, c
->brief
, ss
); continue;
1862 rv
|= ExamineBadblocks(dv
->devname
, c
->brief
, ss
); continue;
1865 rv
|= Wait(dv
->devname
); continue;
1867 rv
|= WaitClean(dv
->devname
, -1, c
->verbose
); continue;
1869 rv
|= Kill_subarray(dv
->devname
, c
->subarray
, c
->verbose
);
1871 case UpdateSubarray
:
1872 if (c
->update
== NULL
) {
1873 pr_err("-U/--update must be specified with --update-subarray\n");
1877 rv
|= Update_subarray(dv
->devname
, c
->subarray
,
1878 c
->update
, ident
, c
->verbose
);
1881 rv
|= Dump_metadata(dv
->devname
, dump_directory
, c
, ss
);
1884 rv
|= Restore_metadata(dv
->devname
, dump_directory
, c
, ss
,
1885 (dv
== devlist
&& dv
->next
== NULL
));
1888 rv
|= SetAction(dv
->devname
, c
->action
);
1891 if (dv
->devname
[0] == '/')
1892 mdfd
= open_mddev(dv
->devname
, 1);
1894 mdfd
= open_dev(dv
->devname
);
1896 pr_err("Cannot open %s\n", dv
->devname
);
1899 switch(dv
->disposition
) {
1902 rv
|= Manage_run(dv
->devname
, mdfd
, c
); break;
1904 rv
|= Manage_stop(dv
->devname
, mdfd
, c
->verbose
, 0); break;
1906 rv
|= Manage_ro(dv
->devname
, mdfd
, 1); break;
1908 rv
|= Manage_ro(dv
->devname
, mdfd
, -1); break;
1917 int SetAction(char *dev
, char *action
)
1919 int fd
= open(dev
, O_RDONLY
);
1922 pr_err("Couldn't open %s: %s\n", dev
, strerror(errno
));
1925 sysfs_init(&mdi
, fd
, NULL
);
1927 if (!mdi
.sys_name
[0]) {
1928 pr_err("%s is no an md array\n", dev
);
1932 if (sysfs_set_str(&mdi
, NULL
, "sync_action", action
) < 0) {
1933 pr_err("Count not set action for %s to %s: %s\n",
1934 dev
, action
, strerror(errno
));